{
  "data": {
    "a": {
      "slug": "baseten",
      "name": "Baseten",
      "vendor": "Baseten",
      "vendorUrl": "https://www.baseten.co",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Dedicated model deployments packaged with the open-source Truss framework and served behind a per-model HTTPS endpoint, with autoscaling from zero replicas, async inference, a management API and per-minute GPU billing from T4 to B200.",
      "url": "https://www.anchorterminal.com/tools/baseten",
      "markdownUrl": "https://www.anchorterminal.com/tools/baseten.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/baseten.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/baseten.json",
      "repo": "https://github.com/basetenlabs/truss",
      "license": "MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.baseten.co",
      "packages": [
        {
          "registry": "pypi",
          "name": "truss"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key from the workspace settings, sent as `Authorization: Bearer $BASETEN_API_KEY` (preferred) or the legacy `Authorization: Api-Key` scheme. Keys created from 1 October 2026 carry a `b10_` prefix. Inference goes to model-\u003cid\u003e.api.baseten.co and management calls to api.baseten.co.",
      "pricing": "usage",
      "pricingNotes": "Basic is $0 a month, pay as you go; Pro and Enterprise add volume discounts. Dedicated deployments bill per minute of replica time, including start-up and idle, and nothing at zero replicas. T4 16 GiB $0.01052 a minute (about $0.63 an hour), L4 24 GiB $0.01414 ($0.85), A10G 24 GiB $0.02012 ($1.21), H100 MIG 40 GiB $0.0625 ($3.75), A100 80 GiB $0.06667 ($4.00), H100 80 GiB $0.10833 ($6.50), B200 180 GiB $0.16633 ($9.98). New accounts get a small credit to try the UI. Model APIs bill per token instead (https://www.baseten.co/pricing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 1200,
        "npmWeekly": null,
        "pypiWeekly": 74496,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.baseten.co",
      "llmsTxt": "https://docs.baseten.co/llms.txt",
      "openapi": "https://api.baseten.co/v1/spec",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "python",
        "llms-txt",
        "open-source",
        "async-jobs",
        "webhooks",
        "enterprise"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.5,
        "grade": "B",
        "agentReady": false,
        "rank": 233,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 1,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 65
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -5,
        "negativeNotes": [
          "-5: a GitHub personal access token for `basetenbot`, exposed in a public Harbor image since March 2023, gave admin and push access to Baseten's main product repository, the GitOps repository that drives its clusters, its Homebrew tap and per-customer private repositories. Reported on 2026-07-13, revoked on 2026-07-14, no misuse found, published with Baseten's approval in September 2026. Deducted less because the fix was quick and documented (https://www.strix.ai/blog/baseten-harbor-github-pat-takeover)"
        ],
        "verdict": "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.",
        "bestFor": "Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.",
        "strengths": [
          "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026",
          "Public OpenAPI spec for the management API at api.baseten.co/v1/spec, and llms.txt with Markdown twins",
          "Rate limits published per endpoint with a `retry_after` field on 429",
          "Free starting credits with no payment method needed until they run out",
          "Truss (MIT) keeps the model package portable, with three releases in September 2026"
        ],
        "weaknesses": [
          "H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed",
          "21 status-page incidents between 31 July and 29 September 2026, mostly single-cluster 5xx",
          "A bot token with admin access to the product and GitOps repositories sat exposed from March 2023 until July 2026",
          "No pagination on management list endpoints and no idempotency keys",
          "No SLA below Enterprise, no bug bounty and no security.txt"
        ],
        "agentNotes": [
          "Create a team key with inference-only permission for calling models and keep full-access keys out of the agent",
          "Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute",
          "Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code",
          "Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes",
          "Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.5
          }
        ],
        "editorialScores": {
          "ergonomics": 55,
          "maintenance": 90,
          "payments": 40,
          "reliability": 80,
          "schema": 84,
          "security": 82,
          "transparency": 58
        },
        "provenanceScore": 71
      },
      "connect": {
        "install": "pip install truss",
        "http": "curl -X POST \"https://model-$BASETEN_MODEL_ID.api.baseten.co/environments/production/predict\" \\\n  -H \"Authorization: Bearer $BASETEN_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"prompt\":\"Hello, world!\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/baseten"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GiB",
          "unit": "gpu-hour",
          "usd": 6.5,
          "note": "$0.10833 a minute"
        },
        {
          "item": "B200 180 GiB",
          "unit": "gpu-hour",
          "usd": 9.98,
          "note": "$0.16633 a minute"
        },
        {
          "item": "A100 80 GiB",
          "unit": "gpu-hour",
          "usd": 4,
          "note": "$0.06667 a minute"
        },
        {
          "item": "H100 MIG 40 GiB",
          "unit": "gpu-hour",
          "usd": 3.75,
          "note": "$0.0625 a minute"
        },
        {
          "item": "A10G 24 GiB",
          "unit": "gpu-hour",
          "usd": 1.21,
          "note": "$0.02012 a minute"
        },
        {
          "item": "L4 24 GiB",
          "unit": "gpu-hour",
          "usd": 0.85,
          "note": "$0.01414 a minute"
        },
        {
          "item": "T4 16 GiB",
          "unit": "gpu-hour",
          "usd": 0.63,
          "note": "$0.01052 a minute"
        }
      ],
      "provenance": {
        "legalEntity": "Baseten Labs, Inc.",
        "domain": "baseten.co",
        "domainRegistered": "",
        "endpointOnVendorDomain": true,
        "terms": "https://www.baseten.co/terms-and-conditions/",
        "privacy": "https://www.baseten.co/privacy-policy/",
        "statusPage": "https://status.baseten.co",
        "changelog": "https://www.baseten.co/changelog/",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms name Baseten Labs, Inc. under California law with venue in San Francisco. The privacy policy gives 560 Davis St., Suite 250, San Francisco.",
          "Inference runs on model-\u003cid\u003e.api.baseten.co, a subdomain of the vendor domain.",
          "www.baseten.co/.well-known/security.txt returns 404.",
          "The .co registry's RDAP server couldn't be reached, so the registration date is unrecorded."
        ],
        "score": 71
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/baseten.json",
      "live": {
        "slug": "baseten",
        "probe": {
          "target": "https://api.baseten.co",
          "method": "get",
          "lastAt": "2026-10-08T19:52:45.872637881Z",
          "lastOk": true,
          "lastStatus": 202,
          "lastMs": 454,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 460,
          "p95ms24h": 519,
          "samples24h": 272,
          "samples30d": 1941,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 225,
              "ok": 225
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.baseten.co",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-08T19:50:26.050899126Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "basetenlabs/truss",
            "version": "v0.18.33",
            "released": "2026-10-06",
            "seenAt": "2026-10-08T16:02:10.626982738Z"
          },
          {
            "registry": "pypi",
            "name": "truss",
            "version": "0.18.33",
            "released": "2026-10-06",
            "seenAt": "2026-10-08T16:02:06.626291704Z"
          }
        ],
        "githubStars": 1214,
        "pypiWeekly": 62737,
        "securityTxt": {
          "url": "https://baseten.co/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:37.892521999Z"
        },
        "llmsTxt": {
          "url": "https://docs.baseten.co/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:06.558742297Z"
        },
        "domain": {
          "domain": "baseten.co",
          "checkedAt": "2026-10-04T13:09:05.701625104Z"
        },
        "pages": [
          {
            "url": "https://www.baseten.co/changelog/",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:31.551352111Z",
            "changedAt": "2026-10-06T16:14:05.469135663Z",
            "fingerprint": "fcc7d4bab97d"
          },
          {
            "url": "https://www.baseten.co/pricing/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:33.809519828Z",
            "changedAt": "2026-10-06T16:14:07.823033015Z",
            "fingerprint": "e41ad7119a6b"
          },
          {
            "url": "https://www.baseten.co/privacy-policy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:35.678713031Z",
            "changedAt": "2026-10-06T16:14:09.752979771Z",
            "fingerprint": "09df4e188f65"
          },
          {
            "url": "https://www.baseten.co/terms-and-conditions/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:38.014111749Z",
            "changedAt": "2026-10-06T16:14:11.853258349Z",
            "fingerprint": "a8b301619021"
          }
        ],
        "updatedAt": "2026-10-08T19:52:45.872637881Z"
      }
    },
    "answer": "Baseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category.",
    "b": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 512,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:52:47.839655439Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 248,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 252,
          "p95ms24h": 326,
          "samples24h": 27,
          "samples30d": 27,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 27,
              "ok": 27
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:38:18.718586754Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:52:47.839655439Z"
      }
    },
    "facts": [
      {
        "a": "HTTP API",
        "b": "Model platform",
        "name": "Kind"
      },
      {
        "a": "Baseten",
        "b": "Cerebrium Inc.",
        "name": "Vendor"
      },
      {
        "a": "https://api.baseten.co",
        "b": "https://rest.cerebrium.ai",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT",
        "b": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-09-16",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "no date given",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "no date given",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "1.2k stars, 74k PyPI/wk",
        "b": "920 PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "3.5/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Baseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category.",
        "question": "Which is better for AI agents, Baseten or Cerebrium?"
      },
      {
        "answer": "Yes. Baseten has a hosted endpoint at https://api.baseten.co and Cerebrium at https://rest.cerebrium.ai.",
        "question": "Can an agent call Baseten and Cerebrium without installing anything?"
      },
      {
        "answer": "Baseten is open source (MIT). No open-source release is listed for Cerebrium.",
        "question": "Are Baseten and Cerebrium open source?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 80 against 48",
          "Schema \u0026 documentation, 84 against 70",
          "Agent ergonomics, 55 against 49",
          "Security \u0026 auth, 82 against 60",
          "Payments \u0026 pricing, 40 against 30",
          "Maintenance \u0026 community, 90 against 75"
        ],
        "also": [
          "Open source"
        ],
        "goodFor": "Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.",
        "slug": "baseten",
        "watchFor": "H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed"
      },
      {
        "aheadOn": null,
        "also": [
          "No incidents deducted, where Baseten loses 5 points for them"
        ],
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-beam.json",
        "title": "Baseten vs Beam",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-beam"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-coreweave.json",
        "title": "Baseten vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-koyeb.json",
        "title": "Baseten vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-lambda.json",
        "title": "Baseten vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-modal.json",
        "title": "Baseten vs Modal",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-northflank.json",
        "title": "Baseten vs Northflank",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.json",
        "title": "Baseten vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-runpod.json",
        "title": "Baseten vs Runpod",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai.json",
        "title": "Baseten vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "baseten": 80,
        "by": 32,
        "cerebrium": 48,
        "edge": "baseten",
        "key": "reliability",
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "baseten": 84,
        "by": 14,
        "cerebrium": 70,
        "edge": "baseten",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "baseten": 55,
        "by": 6,
        "cerebrium": 49,
        "edge": "baseten",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "baseten": 82,
        "by": 22,
        "cerebrium": 60,
        "edge": "baseten",
        "key": "security",
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "baseten": 40,
        "by": 10,
        "cerebrium": 30,
        "edge": "baseten",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "baseten": 90,
        "by": 15,
        "cerebrium": 75,
        "edge": "baseten",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "baseten": 65,
        "by": 2,
        "cerebrium": 63,
        "edge": "baseten",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Baseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category. Both do compute gpu.",
    "verdicts": {
      "baseten": "Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.",
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium",
    "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md",
    "slim": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.min.md"
  },
  "markdown": "Baseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category. Both do compute gpu.\n\n- Baseten: grade B, 66.5/100, rank #233 of 722. Markdown https://www.anchorterminal.com/tools/baseten.md · JSON https://www.anchorterminal.com/api/v1/tools/baseten.json\n- Cerebrium: grade C, 55.3/100, rank #512 of 722. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n\n## Which one, for what\n\n### Baseten (B)\n\nGood for: Teams that want one model behind a production endpoint with real autoscaling knobs, environments and scoped keys.\n\nAhead on:\n- Reliability, 80 against 48\n- Schema \u0026 documentation, 84 against 70\n- Agent ergonomics, 55 against 49\n- Security \u0026 auth, 82 against 60\n- Payments \u0026 pricing, 40 against 30\n- Maintenance \u0026 community, 90 against 75\n\nAlso in its favour:\n- Open source\n\nWatch for: H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAlso in its favour:\n- No incidents deducted, where Baseten loses 5 points for them\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n\n## Score by category\n\n| Category | Weight | Baseten | Cerebrium | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 80 | 48 | Baseten +32 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 84 | 70 | Baseten +14 |\n| Agent ergonomics | 13% (16.2 this run) | 55 | 49 | Baseten +6 |\n| Security \u0026 auth | 14% (17.5 this run) | 82 | 60 | Baseten +22 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 30 | Baseten +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 90 | 75 | Baseten +15 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 65 | 63 | Baseten +2 |\n| Negative events | ≤15 | -5 | 0 | |\n| **Total** | | **66.5 · B** | **55.3 · C** | |\n\n## Facts side by side\n\n| Fact | Baseten | Cerebrium |\n| --- | --- | --- |\n| Kind | HTTP API | Model platform |\n| Vendor | Baseten | Cerebrium Inc. |\n| Hosted endpoint | `https://api.baseten.co` | `https://rest.cerebrium.ai` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | MIT | Proprietary service under Cerebrium's terms of service. The CLI is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-28 | 2026-09-16 |\n| Terms last updated | no date given | no date given |\n| Privacy policy last updated | no date given | no date given |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | not found in the text | yes |\n| Terms restrict benchmarking | yes | not found in the text |\n| Terms or service can change without notice | not found in the text | yes |\n| Arbitration or class-action waiver | not found in the text | not found in the text |\n| Popularity | 1.2k stars, 74k PyPI/wk | 920 PyPI/wk |\n| Agent reviews | 3.5/5 (2) | none |\n\n## Verdicts\n\n**Baseten.** Team API keys scoped to inference-only, metrics-only or a single environment or model, plus a Viewer role since 1 September 2026. H100 at $6.50 and A100 at $4.00 an hour, and start-up and idle replica time are billed.\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n## Before you call either\n\n### Baseten\n\n1. Create a team key with inference-only permission for calling models and keep full-access keys out of the agent\n2. Sleep for `retry_after` seconds on a 429 from api.baseten.co; the activate and deactivate endpoints allow 20 calls a minute\n3. Retry 429, 503 and 529 with backoff, but treat 500 as a bug in your model code\n4. Set `scale_down_delay` below the 900-second default or every burst bills 15 idle minutes\n5. Send payloads over 256 KiB to `/predict`, not `/async_predict`, unless support has raised the async limit\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n## Questions\n\n### Which is better for AI agents, Baseten or Cerebrium?\n\nBaseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category.\n\n### Can an agent call Baseten and Cerebrium without installing anything?\n\nYes. Baseten has a hosted endpoint at https://api.baseten.co and Cerebrium at https://rest.cerebrium.ai.\n\n### Are Baseten and Cerebrium open source?\n\nBaseten is open source (MIT). No open-source release is listed for Cerebrium.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json, and with the fewest tokens: https://www.anchorterminal.com/compare/baseten-vs-cerebrium.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"baseten\", \"b\": \"cerebrium\"}`. From a terminal: `anchor compare baseten cerebrium`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/baseten.json and https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n\n## Other comparisons with Baseten or Cerebrium\n\n- [Baseten vs Beam](https://www.anchorterminal.com/compare/baseten-vs-beam.md)\n- [Baseten vs CoreWeave](https://www.anchorterminal.com/compare/baseten-vs-coreweave.md)\n- [Baseten vs Koyeb](https://www.anchorterminal.com/compare/baseten-vs-koyeb.md)\n- [Baseten vs Lambda Cloud](https://www.anchorterminal.com/compare/baseten-vs-lambda.md)\n- [Baseten vs Modal](https://www.anchorterminal.com/compare/baseten-vs-modal.md)\n- [Baseten vs Northflank](https://www.anchorterminal.com/compare/baseten-vs-northflank.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Baseten vs Runpod](https://www.anchorterminal.com/compare/baseten-vs-runpod.md)\n- [Baseten vs Vast.ai](https://www.anchorterminal.com/compare/baseten-vs-vast-ai.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Baseten vs Cerebrium",
        "url": ""
      }
    ],
    "description": "Baseten scores 66.5 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in every scored category. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Baseten B 66.5",
      "Cerebrium C 55.3",
      "scores"
    ],
    "h1": "Baseten vs Cerebrium",
    "image": "https://www.anchorterminal.com/assets/og/compare-baseten-vs-cerebrium.png",
    "path": "/compare/baseten-vs-cerebrium",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Baseten vs Cerebrium for AI agents, B 66.5 vs C 55.3 | Anchor Terminal",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
  },
  "tokens": {
    "markdown": 2150,
    "slim": 630
  },
  "version": 1
}
