{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 454,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:08:42.947721656Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 326,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 251,
          "p95ms24h": 326,
          "samples24h": 19,
          "samples30d": 19,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 19,
              "ok": 19
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:50:27.054213755Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:08:42.947721656Z"
      }
    },
    "answer": "Cerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community.",
    "b": {
      "slug": "runpod",
      "name": "Runpod",
      "vendor": "Runpod",
      "vendorUrl": "https://www.runpod.io",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Serverless GPU endpoints, queue-based or load-balanced, and rented GPU Pods, billed per second from prepaid credit.",
      "url": "https://www.anchorterminal.com/tools/runpod",
      "markdownUrl": "https://www.anchorterminal.com/tools/runpod.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/runpod.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/runpod.json",
      "repo": "https://github.com/runpod/runpod-python",
      "license": "MIT",
      "transports": [
        "http",
        "streamable-http",
        "stdio"
      ],
      "remoteUrl": "https://api.runpod.ai/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "runpod"
        },
        {
          "registry": "npm",
          "name": "runpod-sdk"
        },
        {
          "registry": "npm",
          "name": "@runpod/mcp-server"
        },
        {
          "registry": "pypi",
          "name": "runpod-flash"
        }
      ],
      "auth": "mixed",
      "authNotes": "API key from the console sent as `Authorization: Bearer` to serverless endpoints at api.runpod.ai/v2/\u003cendpoint-id\u003e and to the management REST API v2 at api.runpod.io/v2. Keys can be All, Read Only or Restricted per serverless endpoint. The deprecated GraphQL API takes the key as `?api_key=` in the URL. The hosted MCP server at mcp.getrunpod.io signs in with OAuth (Sign in with Runpod) or takes an API key as a Bearer header; the local `@runpod/mcp-server` reads `RUNPOD_API_KEY`.",
      "pricing": "usage",
      "pricingNotes": "Prepaid credit, billed per second and rounded up to the nearest second, with no data transfer fees and no free tier. Serverless flex workers an hour are 16 GB A4000 class $0.58, L4, A5000 or RTX 3090 $0.69, RTX 4090 $1.10, RTX Pro 4500 $1.15, A6000 or A40 $1.22, RTX 5090 $1.58, L40, L40S or RTX 6000 Ada $1.75, A100 80 GB $2.72, RTX Pro 6000 $3.49, H100 $4.79, H200 $5.93, B200 $8.64, B300 $9.98. Active (always-on) workers are discounted through sales. Pods run from $0.27 an hour (RTX A5000) to $7.89 (B300) on Community or Secure Cloud. Container disk $0.10 a GB-month, volume disk $0.10 running and $0.20 idle, network volumes $0.07 a GB-month under 1 TB and $0.05 above, high-performance $0.14. Default spend cap $80 an hour. Cards, crypto after KYC, prepaid cards at $100 or more a transaction, invoicing above $5,000 (https://www.runpod.io/pricing, https://docs.runpod.io/serverless/pricing, https://docs.runpod.io/accounts-billing/billing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 314,
        "npmWeekly": 21838,
        "pypiWeekly": 146926,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.runpod.io",
      "llmsTxt": "https://docs.runpod.io/llms.txt",
      "openapi": "https://api.runpod.io/v2/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "prepaid",
        "mcp",
        "oauth",
        "llms-txt",
        "python",
        "typescript",
        "async-jobs",
        "batch",
        "webhooks",
        "enterprise"
      ],
      "lastRelease": "2026-09-15",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 53.5,
        "grade": "D",
        "agentReady": false,
        "rank": 479,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 9,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 47,
          "maintenance": 82,
          "payments": 20,
          "reliability": 35,
          "schema": 81,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Per-second billing across more than a dozen serverless GPU classes, H100 at $4.79 and A100 80 GB at $2.72 an hour. Data-centre outages of 6 to 24 hours in each of July, August and September 2026.",
        "bestFor": "Cost-sensitive inference and batch work that wants the widest GPU choice, from consumer cards to B300, with an MCP control plane.",
        "strengths": [
          "Per-second billing across more than a dozen serverless GPU classes, H100 at $4.79 and A100 80 GB at $2.72 an hour",
          "API keys can be Read Only or restricted per serverless endpoint",
          "OpenAPI file for REST v2, llms.txt and dated release notes with deprecation dates",
          "Official MCP server, hosted with OAuth or local over stdio",
          "Valid security.txt and SOC 2 Type 2 and ISO 27001 in the trust centre"
        ],
        "weaknesses": [
          "Data-centre outages of 6 to 24 hours in each of July, August and September 2026",
          "No published rate limits, 429 guidance, error codes or SLA",
          "The GraphQL API, live until early 2027, takes the API key in the URL",
          "No free tier, prepaid credit only",
          "Three APIs in flight, with REST v1 retiring on 15 November 2026"
        ],
        "agentNotes": [
          "Create a Restricted or Read Only key per endpoint for the agent, not an All key",
          "Use REST v2 at api.runpod.io/v2 with a Bearer header; avoid GraphQL, which puts the key in the URL",
          "Fetch `/run` results within 30 minutes and `/runsync` results within 1 minute, or they're gone",
          "Call `/retry` on a failed job ID rather than submitting a duplicate job",
          "Check `/health` before relying on an endpoint idle for a week, since max workers drop to 0"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 53.5
          }
        ],
        "editorialScores": {
          "ergonomics": 47,
          "maintenance": 82,
          "payments": 20,
          "reliability": 35,
          "schema": 81,
          "security": 60,
          "transparency": 60
        },
        "provenanceScore": 66
      },
      "connect": {
        "install": "pip install runpod  # or npm i runpod-sdk",
        "http": "curl -X POST \"https://api.runpod.ai/v2/$RUNPOD_ENDPOINT_ID/runsync\" \\\n  -H \"Authorization: Bearer $RUNPOD_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"input\":{\"prompt\":\"Hello, world!\"}}'",
        "claudeCode": "claude mcp add --transport http runpod https://mcp.getrunpod.io/ --header \"Authorization: Bearer $RUNPOD_API_KEY\"",
        "config": {
          "mcpServers": {
            "runpod": {
              "args": [
                "-y",
                "@runpod/mcp-server@latest"
              ],
              "command": "npx",
              "env": {
                "RUNPOD_API_KEY": "${RUNPOD_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/runpod"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 4.79,
          "note": "Billed per second"
        },
        {
          "item": "H200 141 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 5.93
        },
        {
          "item": "B200 180 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 8.64
        },
        {
          "item": "A100 80 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 2.72
        },
        {
          "item": "L40S 48 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 1.75
        },
        {
          "item": "RTX 4090 24 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 1.1
        },
        {
          "item": "L4 24 GB serverless flex",
          "unit": "gpu-hour",
          "usd": 0.69
        },
        {
          "item": "Container disk",
          "unit": "gb-month",
          "usd": 0.1
        },
        {
          "item": "Network volume under 1 TB",
          "unit": "gb-month",
          "usd": 0.07,
          "note": "$0.05 above 1 TB, $0.14 high-performance"
        }
      ],
      "provenance": {
        "legalEntity": "Runpod, Inc.",
        "domain": "runpod.io",
        "domainRegistered": "",
        "endpointOnVendorDomain": false,
        "terms": "https://www.runpod.io/legal/terms-of-service",
        "privacy": "https://www.runpod.io/legal/privacy-policy",
        "statusPage": "https://uptime.runpod.io",
        "changelog": "https://docs.runpod.io/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "notes": [
          "Terms effective 24 March 2026 name Runpod, Inc. under Delaware law. The privacy policy gives 329 Bryant St #4D, San Francisco.",
          "Serverless endpoints are served from api.runpod.ai and the management API from api.runpod.io, while the hosted MCP server sits on getrunpod.io.",
          "security.txt points to trust.runpod.io and expires 2027-01-31.",
          "The .io registry's RDAP server rate-limited our lookup, so the registration date is unrecorded."
        ],
        "score": 66
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/runpod.json",
      "live": {
        "slug": "runpod",
        "probe": {
          "target": "https://api.runpod.ai/v2",
          "method": "get",
          "lastAt": "2026-10-08T19:08:57.623696013Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 77,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 54,
          "p95ms24h": 122,
          "samples24h": 272,
          "samples30d": 1933,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 217,
              "ok": 217
            }
          ]
        },
        "vendorStatus": {
          "page": "https://uptime.runpod.io",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:51:10.831410075Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "runpod/runpod-python",
            "version": "v1.12.0",
            "released": "2026-08-10",
            "seenAt": "2026-10-08T16:27:52.721885319Z"
          },
          {
            "registry": "npm",
            "name": "@runpod/mcp-server",
            "version": "4.0.0",
            "seenAt": "2026-10-08T16:27:51.181072709Z"
          },
          {
            "registry": "npm",
            "name": "runpod-sdk",
            "version": "1.1.2",
            "seenAt": "2026-10-08T16:27:50.332110359Z"
          },
          {
            "registry": "pypi",
            "name": "runpod",
            "version": "1.12.0",
            "released": "2026-08-10",
            "seenAt": "2026-10-08T16:27:50.147142378Z"
          },
          {
            "registry": "pypi",
            "name": "runpod-flash",
            "version": "1.20.0",
            "released": "2026-09-24",
            "seenAt": "2026-10-08T16:27:52.644090717Z"
          }
        ],
        "githubStars": 314,
        "npmWeekly": 30649,
        "pypiWeekly": 163037,
        "securityTxt": {
          "url": "https://runpod.io/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-01-31T23:59:59Z",
          "checkedAt": "2026-10-08T15:38:36.644714031Z"
        },
        "llmsTxt": {
          "url": "https://docs.runpod.io/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:51.140150765Z"
        },
        "domain": {
          "domain": "runpod.io",
          "checkedAt": "2026-10-04T13:06:54.960688955Z"
        },
        "pages": [
          {
            "url": "https://docs.runpod.io/release-notes",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-08T18:19:23.565650782Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "64302af06aec"
          },
          {
            "url": "https://docs.runpod.io/serverless/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:19:25.792000525Z",
            "changedAt": "2026-10-06T16:08:35.000217664Z",
            "fingerprint": "e5a52588b7ab"
          },
          {
            "url": "https://www.runpod.io/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:09.061455266Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "4e4c891c31f7"
          },
          {
            "url": "https://www.runpod.io/legal/privacy-policy",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:05.013110293Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "00b9defcdf3d"
          },
          {
            "url": "https://www.runpod.io/legal/terms-of-service",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:07.057529612Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "4163df004c48"
          }
        ],
        "updatedAt": "2026-10-08T19:08:57.623696013Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Runpod",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "https://api.runpod.ai/v2",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, Streamable HTTP, stdio",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "OAuth or key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "2026-09-15",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2026-03-24",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2025-08-07",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "314 stars, 22k npm/wk, 147k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Cerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community.",
        "question": "Which is better for AI agents, Cerebrium or Runpod?"
      },
      {
        "answer": "Yes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Runpod at https://api.runpod.ai/v2.",
        "question": "Can an agent call Cerebrium and Runpod without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 48 against 35",
          "Payments \u0026 pricing, 30 against 20"
        ],
        "also": null,
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 81 against 70",
          "Maintenance \u0026 community, 82 against 75"
        ],
        "also": [
          "Runs on your own machine"
        ],
        "goodFor": "Cost-sensitive inference and batch work that wants the widest GPU choice, from consumer cards to B300, with an MCP control plane.",
        "slug": "runpod",
        "watchFor": "Data-centre outages of 6 to 24 hours in each of July, August and September 2026"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-runpod.json",
        "title": "Baseten vs Runpod",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-runpod.json",
        "title": "Beam vs Runpod",
        "url": "https://www.anchorterminal.com/compare/beam-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-runpod.json",
        "title": "CoreWeave vs Runpod",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-runpod.json",
        "title": "Koyeb vs Runpod",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-runpod.json",
        "title": "Lambda Cloud vs Runpod",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-runpod.json",
        "title": "Modal vs Runpod",
        "url": "https://www.anchorterminal.com/compare/modal-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/northflank-vs-runpod.json",
        "title": "Northflank vs Runpod",
        "url": "https://www.anchorterminal.com/compare/northflank-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.json",
        "title": "Replicate Deployments vs Runpod",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/runpod-vs-vast-ai.json",
        "title": "Runpod vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/runpod-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 13,
        "cerebrium": 48,
        "edge": "cerebrium",
        "key": "reliability",
        "name": "Reliability",
        "runpod": 35,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 11,
        "cerebrium": 70,
        "edge": "runpod",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "runpod": 81,
        "weight": 13
      },
      {
        "by": 2,
        "cerebrium": 49,
        "edge": "cerebrium",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "runpod": 47,
        "weight": 13
      },
      {
        "by": 0,
        "cerebrium": 60,
        "edge": "",
        "key": "security",
        "name": "Security \u0026 auth",
        "runpod": 60,
        "weight": 14
      },
      {
        "by": 10,
        "cerebrium": 30,
        "edge": "cerebrium",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "runpod": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "cerebrium": 75,
        "edge": "runpod",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "runpod": 82,
        "weight": 7
      },
      {
        "by": 0,
        "cerebrium": 63,
        "edge": "",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "runpod": 63,
        "weight": 7
      }
    ],
    "summary": "Cerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "runpod": "Per-second billing across more than a dozen serverless GPU classes, H100 at $4.79 and A100 80 GB at $2.72 an hour. Data-centre outages of 6 to 24 hours in each of July, August and September 2026."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.min.md"
  },
  "markdown": "Cerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #454 of 629. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Runpod: grade D, 53.5/100, rank #479 of 629. Markdown https://www.anchorterminal.com/tools/runpod.md · JSON https://www.anchorterminal.com/api/v1/tools/runpod.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAhead on:\n- Reliability, 48 against 35\n- Payments \u0026 pricing, 30 against 20\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Runpod (D)\n\nGood for: Cost-sensitive inference and batch work that wants the widest GPU choice, from consumer cards to B300, with an MCP control plane.\n\nAhead on:\n- Schema \u0026 documentation, 81 against 70\n- Maintenance \u0026 community, 82 against 75\n\nAlso in its favour:\n- Runs on your own machine\n\nWatch for: Data-centre outages of 6 to 24 hours in each of July, August and September 2026\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Runpod | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 35 | Cerebrium +13 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 81 | Runpod +11 |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 47 | Cerebrium +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 60 | even |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 20 | Cerebrium +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 82 | Runpod +7 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 63 | even |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **53.5 · D** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Runpod |\n| --- | --- | --- |\n| Kind | Model platform | HTTP API |\n| Vendor | Cerebrium Inc. | Runpod |\n| Hosted endpoint | `https://rest.cerebrium.ai` | `https://api.runpod.ai/v2` |\n| Transports | HTTP | HTTP, Streamable HTTP, stdio |\n| Auth | API key | OAuth or key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-16 | 2026-09-15 |\n| Terms last updated | no date given | 2026-03-24 |\n| Privacy policy last updated | no date given | 2025-08-07 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | yes |\n| Terms restrict benchmarking | not found in the text | yes |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | not found in the text | yes |\n| Popularity | 920 PyPI/wk | 314 stars, 22k npm/wk, 147k PyPI/wk |\n| Agent reviews | none | 3/5 (2) |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Runpod.** Per-second billing across more than a dozen serverless GPU classes, H100 at $4.79 and A100 80 GB at $2.72 an hour. Data-centre outages of 6 to 24 hours in each of July, August and September 2026.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Runpod\n\n1. Create a Restricted or Read Only key per endpoint for the agent, not an All key\n2. Use REST v2 at api.runpod.io/v2 with a Bearer header; avoid GraphQL, which puts the key in the URL\n3. Fetch `/run` results within 30 minutes and `/runsync` results within 1 minute, or they're gone\n4. Call `/retry` on a failed job ID rather than submitting a duplicate job\n5. Check `/health` before relying on an endpoint idle for a week, since max workers drop to 0\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Runpod?\n\nCerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community.\n\n### Can an agent call Cerebrium and Runpod without installing anything?\n\nYes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Runpod at https://api.runpod.ai/v2.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-runpod.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"runpod\"}`. From a terminal: `anchor compare cerebrium runpod`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/runpod.json\n\n## Other comparisons with Cerebrium or Runpod\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Runpod](https://www.anchorterminal.com/compare/baseten-vs-runpod.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Runpod](https://www.anchorterminal.com/compare/beam-vs-runpod.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n- [CoreWeave vs Runpod](https://www.anchorterminal.com/compare/coreweave-vs-runpod.md)\n- [Koyeb vs Runpod](https://www.anchorterminal.com/compare/koyeb-vs-runpod.md)\n- [Lambda Cloud vs Runpod](https://www.anchorterminal.com/compare/lambda-vs-runpod.md)\n- [Modal vs Runpod](https://www.anchorterminal.com/compare/modal-vs-runpod.md)\n- [Northflank vs Runpod](https://www.anchorterminal.com/compare/northflank-vs-runpod.md)\n- [Replicate Deployments vs Runpod](https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.md)\n- [Runpod vs Vast.ai](https://www.anchorterminal.com/compare/runpod-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Runpod",
        "url": ""
      }
    ],
    "description": "Cerebrium scores 55.3 (C) on agent readiness against Runpod's 53.5 (D), and leads in 3 of 7 scored categories. Runpod leads on schema \u0026 documentation and maintenance \u0026 community. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Runpod D 53.5",
      "scores"
    ],
    "h1": "Cerebrium vs Runpod",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-runpod.png",
    "path": "/compare/cerebrium-vs-runpod",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Runpod for AI agents, C 55.3 vs D 53.5 | Anchor Terminal",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
  },
  "tokens": {
    "markdown": 2050,
    "slim": 680
  },
  "version": 1
}
