{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 512,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:52:47.839655439Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 248,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 252,
          "p95ms24h": 326,
          "samples24h": 27,
          "samples30d": 27,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 27,
              "ok": 27
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:38:18.718586754Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:52:47.839655439Z"
      }
    },
    "answer": "Modal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories.",
    "b": {
      "slug": "modal",
      "name": "Modal",
      "vendor": "Modal",
      "vendorUrl": "https://modal.com",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Serverless functions, web endpoints, servers and GPU jobs from a Python decorator, with JavaScript and Go SDKs.",
      "url": "https://www.anchorterminal.com/tools/modal",
      "markdownUrl": "https://www.anchorterminal.com/tools/modal.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/modal.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/modal.json",
      "repo": "https://github.com/modal-labs/modal-client",
      "license": "Apache-2.0",
      "transports": [],
      "packages": [
        {
          "registry": "pypi",
          "name": "modal"
        },
        {
          "registry": "npm",
          "name": "modal"
        }
      ],
      "auth": "api-key",
      "authNotes": "No public REST API for deploying. The SDKs and CLI authenticate with a token ID and secret from `modal token new`, read from `MODAL_TOKEN_ID` and `MODAL_TOKEN_SECRET` or `~/.modal.toml`; tokens can carry a TTL. Deployed web endpoints are open by default and can be locked with proxy tokens sent as `Modal-Key` and `Modal-Secret` headers. Servers and Endpoints require a proxy token by default, sent as `Authorization: Bearer \u003cid\u003e.\u003csecret\u003e`.",
      "pricing": "freemium",
      "pricingNotes": "Starter is $0 a month with $30 of compute included every month, 3 seats, 100 containers and 10 concurrent GPUs. Team is $250 a month plus compute with $100 included, unlimited seats, 5,000 containers and 50 concurrent GPUs. Enterprise is custom. GPUs bill per second with nothing charged at zero containers. T4 $0.000164, L4 $0.000222, A10 $0.000306, L40S $0.000542, A100 40 GB $0.000583, A100 80 GB $0.000694, RTX PRO 6000 $0.000842, H100 $0.001097, H200 $0.001261, B200 $0.001736 and B300 $0.001972 a second. The pricing page lists CPU at $0.0000131 a core-second (0.125 core minimum per container) and memory at $0.00000222 a GiB-second, and volumes at $0.09 a GiB-month after 1 TiB free (https://modal.com/pricing).",
      "priceSummary": "$250 / mo",
      "where": "local",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 514,
        "npmWeekly": 940973,
        "pypiWeekly": 10146778,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://modal.com/docs/guide",
      "llmsTxt": "https://modal.com/llms.txt",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "no-card",
        "python",
        "typescript",
        "go",
        "llms-txt",
        "enterprise"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63.6,
        "grade": "B",
        "agentReady": false,
        "rank": 302,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 2,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 57,
          "maintenance": 85,
          "payments": 30,
          "reliability": 70,
          "schema": 70,
          "security": 68,
          "transparency": 67
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Scale to zero by default, per-second billing and about one-second container boots. No REST API or OpenAPI spec for deploying or invoking Functions.",
        "bestFor": "Python teams that want GPU functions, batch jobs and HTTP endpoints from one decorator with scale to zero.",
        "strengths": [
          "Scale to zero by default, per-second billing and about one-second container boots",
          "Retention stated per data type (inputs and outputs up to 7 days, logs 1 to 30 days)",
          "Python, JavaScript and Go SDKs, with llms.txt and dated release notes",
          "Four short incidents on the status page between July and September 2026",
          "SOC 2 Type 2, a private HackerOne programme and published disclosure response times"
        ],
        "weaknesses": [
          "No REST API or OpenAPI spec for deploying or invoking Functions",
          "Web endpoints are open by default until proxy tokens are added",
          "RBAC, audit logs and HIPAA only on Enterprise",
          "No published SLA, and Starter caps concurrent GPUs at 10",
          "Region pinning costs 1.15 to 1.75 times the base price"
        ],
        "agentNotes": [
          "Create a proxy token and require it on every web endpoint before sharing the URL; endpoints are public by default",
          "Pass a list to `gpu=` (for example `[\"H100\", \"A100-80GB\"]`) so a job still runs when the first choice is unavailable",
          "Set `scaledown_window` and `min_containers` explicitly; the defaults are 60 seconds and 0",
          "Use `.spawn()` and poll the call ID for long work instead of holding a web request open",
          "Keep web endpoint traffic under 200 requests a second or ask Modal to raise the limit"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 4,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63.6
          }
        ],
        "editorialScores": {
          "ergonomics": 57,
          "maintenance": 85,
          "payments": 30,
          "reliability": 70,
          "schema": 70,
          "security": 68,
          "transparency": 63
        },
        "provenanceScore": 71
      },
      "connect": {
        "install": "pip install modal \u0026\u0026 modal setup",
        "http": "curl -X POST \"https://$MODAL_WORKSPACE--my-app-predict.modal.run\" \\\n  -H \"Modal-Key: $MODAL_PROXY_KEY\" -H \"Modal-Secret: $MODAL_PROXY_SECRET\" \\\n  -H \"Content-Type: application/json\" -d '{\"prompt\":\"hello\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/modal"
      },
      "sameCompany": [
        "modal-sandboxes"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.95,
          "note": "$0.001097 a second, may be upgraded to H200 at the same price"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.54,
          "note": "$0.001261 a second"
        },
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.25,
          "note": "$0.001736 a second"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second"
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second"
        },
        {
          "item": "Team plan",
          "unit": "month",
          "usd": 250,
          "note": "Plus compute, $100 included, 50 concurrent GPUs"
        }
      ],
      "provenance": {
        "legalEntity": "Modal Labs, Inc.",
        "domain": "modal.com",
        "domainRegistered": "1999-03-18",
        "domainNote": "modal.com was registered in 1999, long before Modal Labs, so the domain was bought later.",
        "endpointOnVendorDomain": false,
        "terms": "https://modal.com/legal/terms",
        "privacy": "https://modal.com/legal/privacy-policy",
        "statusPage": "https://status.modal.com",
        "changelog": "https://modal.com/docs/sdk/py/releases",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms (May 2026) name Modal Labs, Inc., a Delaware corporation, under California law.",
          "Deployed web endpoints and Servers are served from *.modal.run, a separate domain from modal.com. Deployment itself goes through the SDK, so there's no public API base URL to check.",
          "modal.com/.well-known/security.txt returns 404. The security guide gives security@modal.com and a private HackerOne programme.",
          "Modal Sandboxes are listed separately under code sandboxes."
        ],
        "score": 71
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/modal.json",
      "live": {
        "slug": "modal",
        "vendorStatus": {
          "page": "https://status.modal.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:50:52.712660391Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "modal",
            "version": "0.11.0",
            "seenAt": "2026-10-08T16:21:43.713175523Z"
          },
          {
            "registry": "pypi",
            "name": "modal",
            "version": "1.6.1",
            "released": "2026-10-03",
            "seenAt": "2026-10-08T16:21:43.571967985Z"
          }
        ],
        "githubStars": 522,
        "npmWeekly": 976721,
        "pypiWeekly": 10794735,
        "securityTxt": {
          "url": "https://modal.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:35.387911158Z"
        },
        "llmsTxt": {
          "url": "https://modal.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:41.354164105Z"
        },
        "domain": {
          "domain": "modal.com",
          "registered": "1999-03-18",
          "source": "https://rdap.verisign.com/com/v1/domain/modal.com",
          "checkedAt": "2026-10-04T13:03:51.14581991Z"
        },
        "updatedAt": "2026-10-08T17:50:52.712660391Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "Model platform",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Modal",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "no (local only)",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "2026-09-28",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2026-05-01",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2023-05-17",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "514 stars, 941k npm/wk, 10.1M PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "4/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Modal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories.",
        "question": "Which is better for AI agents, Cerebrium or Modal?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": null,
        "also": [
          "A hosted endpoint, with nothing to install"
        ],
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Reliability, 70 against 48",
          "Agent ergonomics, 57 against 49",
          "Security \u0026 auth, 68 against 60",
          "Maintenance \u0026 community, 85 against 75"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Python teams that want GPU functions, batch jobs and HTTP endpoints from one decorator with scale to zero.",
        "slug": "modal",
        "watchFor": "No REST API or OpenAPI spec for deploying or invoking Functions"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-modal.json",
        "title": "Baseten vs Modal",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-modal.json",
        "title": "Beam vs Modal",
        "url": "https://www.anchorterminal.com/compare/beam-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-modal.json",
        "title": "CoreWeave vs Modal",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-modal.json",
        "title": "Koyeb vs Modal",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-modal.json",
        "title": "Lambda Cloud vs Modal",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-northflank.json",
        "title": "Modal vs Northflank",
        "url": "https://www.anchorterminal.com/compare/modal-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.json",
        "title": "Modal vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-runpod.json",
        "title": "Modal vs Runpod",
        "url": "https://www.anchorterminal.com/compare/modal-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-vast-ai.json",
        "title": "Modal vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/modal-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 22,
        "cerebrium": 48,
        "edge": "modal",
        "key": "reliability",
        "modal": 70,
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 0,
        "cerebrium": 70,
        "edge": "",
        "key": "schema",
        "modal": 70,
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 8,
        "cerebrium": 49,
        "edge": "modal",
        "key": "ergonomics",
        "modal": 57,
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 8,
        "cerebrium": 60,
        "edge": "modal",
        "key": "security",
        "modal": 68,
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 0,
        "cerebrium": 30,
        "edge": "",
        "key": "payments",
        "modal": 30,
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 10,
        "cerebrium": 75,
        "edge": "modal",
        "key": "maintenance",
        "modal": 85,
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 4,
        "cerebrium": 63,
        "edge": "modal",
        "key": "transparency",
        "modal": 67,
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Modal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "modal": "Scale to zero by default, per-second billing and about one-second container boots. No REST API or OpenAPI spec for deploying or invoking Functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-modal",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.min.md"
  },
  "markdown": "Modal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #512 of 722. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Modal: grade B, 63.6/100, rank #302 of 722. Markdown https://www.anchorterminal.com/tools/modal.md · JSON https://www.anchorterminal.com/api/v1/tools/modal.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAlso in its favour:\n- A hosted endpoint, with nothing to install\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Modal (B)\n\nGood for: Python teams that want GPU functions, batch jobs and HTTP endpoints from one decorator with scale to zero.\n\nAhead on:\n- Reliability, 70 against 48\n- Agent ergonomics, 57 against 49\n- Security \u0026 auth, 68 against 60\n- Maintenance \u0026 community, 85 against 75\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: No REST API or OpenAPI spec for deploying or invoking Functions\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Modal | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 70 | Modal +22 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 70 | even |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 57 | Modal +8 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 68 | Modal +8 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 30 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 85 | Modal +10 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 67 | Modal +4 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **63.6 · B** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Modal |\n| --- | --- | --- |\n| Kind | Model platform | Model platform |\n| Vendor | Cerebrium Inc. | Modal |\n| Hosted endpoint | `https://rest.cerebrium.ai` | no (local only) |\n| Transports | HTTP |  |\n| Auth | API key | API key |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-16 | 2026-09-28 |\n| Terms last updated | no date given | 2026-05-01 |\n| Privacy policy last updated | no date given | 2023-05-17 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | not found in the text |\n| Terms restrict benchmarking | not found in the text | not found in the text |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | not found in the text | not found in the text |\n| Popularity | 920 PyPI/wk | 514 stars, 941k npm/wk, 10.1M PyPI/wk |\n| Agent reviews | none | 4/5 (2) |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Modal.** Scale to zero by default, per-second billing and about one-second container boots. No REST API or OpenAPI spec for deploying or invoking Functions.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Modal\n\n1. Create a proxy token and require it on every web endpoint before sharing the URL; endpoints are public by default\n2. Pass a list to `gpu=` (for example `[\"H100\", \"A100-80GB\"]`) so a job still runs when the first choice is unavailable\n3. Set `scaledown_window` and `min_containers` explicitly; the defaults are 60 seconds and 0\n4. Use `.spawn()` and poll the call ID for long work instead of holding a web request open\n5. Keep web endpoint traffic under 200 requests a second or ask Modal to raise the limit\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Modal?\n\nModal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-modal.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-modal.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"modal\"}`. From a terminal: `anchor compare cerebrium modal`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/modal.json\n\n## Other comparisons with Cerebrium or Modal\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Modal](https://www.anchorterminal.com/compare/baseten-vs-modal.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Modal](https://www.anchorterminal.com/compare/beam-vs-modal.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n- [CoreWeave vs Modal](https://www.anchorterminal.com/compare/coreweave-vs-modal.md)\n- [Koyeb vs Modal](https://www.anchorterminal.com/compare/koyeb-vs-modal.md)\n- [Lambda Cloud vs Modal](https://www.anchorterminal.com/compare/lambda-vs-modal.md)\n- [Modal vs Northflank](https://www.anchorterminal.com/compare/modal-vs-northflank.md)\n- [Modal vs Replicate Deployments](https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.md)\n- [Modal vs Runpod](https://www.anchorterminal.com/compare/modal-vs-runpod.md)\n- [Modal vs Vast.ai](https://www.anchorterminal.com/compare/modal-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Modal",
        "url": ""
      }
    ],
    "description": "Modal scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 5 of 7 scored categories. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Modal B 63.6",
      "scores"
    ],
    "h1": "Cerebrium vs Modal",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-modal.png",
    "path": "/compare/cerebrium-vs-modal",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Modal for AI agents, C 55.3 vs B 63.6 | Anchor Terminal",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
  },
  "tokens": {
    "markdown": 2000,
    "slim": 580
  },
  "version": 1
}
