{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 454,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:08:42.947721656Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 326,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 251,
          "p95ms24h": 326,
          "samples24h": 19,
          "samples30d": 19,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 19,
              "ok": 19
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:50:27.054213755Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:08:42.947721656Z"
      }
    },
    "answer": "Cerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics.",
    "b": {
      "slug": "lambda",
      "name": "Lambda Cloud",
      "vendor": "Lambda",
      "vendorUrl": "https://lambda.ai",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "On-demand GPU virtual machines and clusters, with an API for provisioning compute and persistent storage.",
      "url": "https://www.anchorterminal.com/tools/lambda",
      "markdownUrl": "https://www.anchorterminal.com/tools/lambda.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/lambda.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/lambda.json",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://cloud.lambda.ai/api/v1",
      "packages": [],
      "auth": "api-key",
      "authNotes": "API key created at cloud.lambda.ai/api-keys and sent as `Authorization: Bearer`. HTTP Basic with the key as the username (`curl -u 'KEY:'`) still works as a legacy option. SSH keys registered in the account are injected into launched instances.",
      "pricing": "usage",
      "pricingNotes": "On-demand, per GPU an hour, Tesla V100 16 GB $0.79, A100 40 GB $1.99, A100 80 GB $2.79, H100 SXM 80 GB $3.99, B200 180 GB $6.69. 1-Click Clusters of HGX B200 are quoted per GPU-hour at $9.86 for 16 GPUs, $9.36 for 64 and $8.87 for 256 or more on two-week to one-year commitments. Instances bill in one-minute increments from the moment they pass health checks until you terminate them, invoiced weekly. Filesystems bill per GB used a month in one-hour increments. No free tier (https://lambda.ai/pricing, https://docs.lambda.ai/public-cloud/billing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.lambda.ai/public-cloud/on-demand/",
      "openapi": "https://docs.lambda.ai/api/cloud/spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "openapi",
        "enterprise"
      ],
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 50,
        "grade": "D",
        "agentReady": false,
        "rank": 523,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 10,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 58
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.",
        "bestFor": "Agents or pipelines that need a whole GPU machine with SSH for hours, such as training or batch jobs, and will manage the lifecycle themselves.",
        "strengths": [
          "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing",
          "OpenAPI 3.1 spec with documented error codes, a `suggestion` field and cursor pagination",
          "Audit events endpoint filterable by time and resource type",
          "Published rate limits, one request a second and one launch every 12 seconds",
          "Trust portal with SOC 2 Type 2 and four ISO certifications, and a named subprocessor list"
        ],
        "weaknesses": [
          "No scale to zero, autoscaling or endpoints; an idle VM bills until terminated",
          "Ten status incidents in two months, including a two-day regional outage in August 2026",
          "No changelog, no llms.txt and no official SDK",
          "API keys have no scopes and launch has no idempotency key",
          "security.txt expired on 1 June 2026 and no free tier"
        ],
        "agentNotes": [
          "Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch",
          "Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`",
          "Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change",
          "List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine",
          "Terminate the instance in a `finally` block; billing runs by the minute until you do"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 50
          }
        ],
        "editorialScores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 48
        },
        "provenanceScore": 67
      },
      "connect": {
        "http": "curl \"https://cloud.lambda.ai/api/v1/instance-types\" -H \"Authorization: Bearer $LAMBDA_API_KEY\""
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/lambda"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 3.99,
          "note": "Billed per minute"
        },
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.69
        },
        {
          "item": "A100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 2.79
        },
        {
          "item": "A100 SXM 40 GB",
          "unit": "gpu-hour",
          "usd": 1.99
        },
        {
          "item": "Tesla V100 16 GB",
          "unit": "gpu-hour",
          "usd": 0.79
        },
        {
          "item": "HGX B200 1-Click Cluster, 16 GPUs",
          "unit": "gpu-hour",
          "usd": 9.86,
          "note": "Two-week to one-year commitment; $8.87 at 256 GPUs or more"
        }
      ],
      "provenance": {
        "legalEntity": "Lambda, Inc.",
        "domain": "lambda.ai",
        "domainRegistered": "",
        "domainNote": "Lambda moved from lambdalabs.com, registered 2008-05-29, to lambda.ai. The .ai registry's RDAP server rate-limited our lookup of the new domain.",
        "endpointOnVendorDomain": true,
        "terms": "https://lambda.ai/legal/terms-of-service",
        "privacy": "https://lambda.ai/legal/privacy-policy",
        "statusPage": "https://status.lambda.ai",
        "changelog": "",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "notes": [
          "Terms dated August 2025 and the privacy policy of 1 January 2026 name Lambda, Inc., 2510 Zanker Road, San Jose, California.",
          "The API runs on cloud.lambda.ai, a subdomain of the vendor domain.",
          "security.txt expired 2026-06-01 and has no Policy field.",
          "No public changelog found for the cloud. The OpenAPI spec reports version 1.10.0."
        ],
        "score": 67
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/lambda.json",
      "live": {
        "slug": "lambda",
        "probe": {
          "target": "https://cloud.lambda.ai/api/v1",
          "method": "get",
          "lastAt": "2026-10-08T19:08:51.093682836Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 905,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 99.95,
          "p50ms24h": 605,
          "p95ms24h": 885,
          "samples24h": 272,
          "samples30d": 1933,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 271
            },
            {
              "date": "2026-10-08",
              "probes": 217,
              "ok": 217
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.lambda.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-08T19:06:44.684478555Z"
        },
        "securityTxt": {
          "url": "https://lambda.ai/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-06-01T16:00:00Z",
          "checkedAt": "2026-10-08T15:38:48.494823468Z"
        },
        "domain": {
          "domain": "lambda.ai",
          "registered": "2017-12-16",
          "source": "https://rdap.identitydigital.services/rdap/domain/lambda.ai",
          "checkedAt": "2026-10-04T13:04:37.89683091Z"
        },
        "pages": [
          {
            "url": "https://lambda.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:21:12.626677302Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ad7e156559d1"
          },
          {
            "url": "https://lambda.ai/legal/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:21:08.557293205Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "921eec75dd83"
          },
          {
            "url": "https://lambda.ai/legal/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:21:10.609958388Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "40e40aee2e7c"
          }
        ],
        "updatedAt": "2026-10-08T19:08:51.093682836Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Lambda",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "https://cloud.lambda.ai/api/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$9.86 per GPU-hour",
        "name": "Price for compute gpu"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "none",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "no",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "none",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2025-08-01",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2026-01-01",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "none",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Cerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics.",
        "question": "Which is better for AI agents, Cerebrium or Lambda Cloud?"
      },
      {
        "answer": "Yes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Lambda Cloud at https://cloud.lambda.ai/api/v1.",
        "question": "Can an agent call Cerebrium and Lambda Cloud without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Payments \u0026 pricing, 30 against 20",
          "Maintenance \u0026 community, 75 against 5",
          "Transparency \u0026 trust, 63 against 58"
        ],
        "also": null,
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Agent ergonomics, 63 against 49"
        ],
        "also": null,
        "goodFor": "Agents or pipelines that need a whole GPU machine with SSH for hours, such as training or batch jobs, and will manage the lifecycle themselves.",
        "slug": "lambda",
        "watchFor": "No scale to zero, autoscaling or endpoints; an idle VM bills until terminated"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-lambda.json",
        "title": "Baseten vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-lambda.json",
        "title": "Beam vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/beam-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-lambda.json",
        "title": "CoreWeave vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-lambda.json",
        "title": "Koyeb vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-modal.json",
        "title": "Lambda Cloud vs Modal",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-northflank.json",
        "title": "Lambda Cloud vs Northflank",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.json",
        "title": "Lambda Cloud vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-runpod.json",
        "title": "Lambda Cloud vs Runpod",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-vast-ai.json",
        "title": "Lambda Cloud vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 2,
        "cerebrium": 48,
        "edge": "lambda",
        "key": "reliability",
        "lambda": 50,
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 1,
        "cerebrium": 70,
        "edge": "cerebrium",
        "key": "schema",
        "lambda": 69,
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 14,
        "cerebrium": 49,
        "edge": "lambda",
        "key": "ergonomics",
        "lambda": 63,
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 0,
        "cerebrium": 60,
        "edge": "",
        "key": "security",
        "lambda": 60,
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 10,
        "cerebrium": 30,
        "edge": "cerebrium",
        "key": "payments",
        "lambda": 20,
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 70,
        "cerebrium": 75,
        "edge": "cerebrium",
        "key": "maintenance",
        "lambda": 5,
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 5,
        "cerebrium": 63,
        "edge": "cerebrium",
        "key": "transparency",
        "lambda": 58,
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Cerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "lambda": "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.min.md"
  },
  "markdown": "Cerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #454 of 629. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Lambda Cloud: grade D, 50/100, rank #523 of 629. Markdown https://www.anchorterminal.com/tools/lambda.md · JSON https://www.anchorterminal.com/api/v1/tools/lambda.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAhead on:\n- Payments \u0026 pricing, 30 against 20\n- Maintenance \u0026 community, 75 against 5\n- Transparency \u0026 trust, 63 against 58\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Lambda Cloud (D)\n\nGood for: Agents or pipelines that need a whole GPU machine with SSH for hours, such as training or batch jobs, and will manage the lifecycle themselves.\n\nAhead on:\n- Agent ergonomics, 63 against 49\n\nWatch for: No scale to zero, autoscaling or endpoints; an idle VM bills until terminated\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Lambda Cloud | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 50 | Lambda Cloud +2 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 69 | Cerebrium +1 |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 63 | Lambda Cloud +14 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 60 | even |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 20 | Cerebrium +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 5 | Cerebrium +70 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 58 | Cerebrium +5 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **50 · D** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Lambda Cloud |\n| --- | --- | --- |\n| Kind | Model platform | HTTP API |\n| Vendor | Cerebrium Inc. | Lambda |\n| Hosted endpoint | `https://rest.cerebrium.ai` | `https://cloud.lambda.ai/api/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| Price for compute gpu | not published | $9.86 per GPU-hour |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | none |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| Last release | 2026-09-16 | none |\n| Terms last updated | no date given | 2025-08-01 |\n| Privacy policy last updated | no date given | 2026-01-01 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | not found in the text |\n| Terms restrict benchmarking | not found in the text | yes |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | not found in the text | yes |\n| Popularity | 920 PyPI/wk | none |\n| Agent reviews | none | 3/5 (2) |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Lambda Cloud.** H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Lambda Cloud\n\n1. Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch\n2. Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`\n3. Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change\n4. List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine\n5. Terminate the instance in a `finally` block; billing runs by the minute until you do\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Lambda Cloud?\n\nCerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics.\n\n### Can an agent call Cerebrium and Lambda Cloud without installing anything?\n\nYes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Lambda Cloud at https://cloud.lambda.ai/api/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-lambda.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"lambda\"}`. From a terminal: `anchor compare cerebrium lambda`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/lambda.json\n\n## Other comparisons with Cerebrium or Lambda Cloud\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Lambda Cloud](https://www.anchorterminal.com/compare/baseten-vs-lambda.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Lambda Cloud](https://www.anchorterminal.com/compare/beam-vs-lambda.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n- [CoreWeave vs Lambda Cloud](https://www.anchorterminal.com/compare/coreweave-vs-lambda.md)\n- [Koyeb vs Lambda Cloud](https://www.anchorterminal.com/compare/koyeb-vs-lambda.md)\n- [Lambda Cloud vs Modal](https://www.anchorterminal.com/compare/lambda-vs-modal.md)\n- [Lambda Cloud vs Northflank](https://www.anchorterminal.com/compare/lambda-vs-northflank.md)\n- [Lambda Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md)\n- [Lambda Cloud vs Runpod](https://www.anchorterminal.com/compare/lambda-vs-runpod.md)\n- [Lambda Cloud vs Vast.ai](https://www.anchorterminal.com/compare/lambda-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Lambda Cloud",
        "url": ""
      }
    ],
    "description": "Cerebrium scores 55.3 (C) on agent readiness against Lambda Cloud's 50 (D), and leads in 4 of 7 scored categories. Lambda Cloud leads on agent ergonomics. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Lambda Cloud D 50",
      "scores"
    ],
    "h1": "Cerebrium vs Lambda Cloud",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-lambda.png",
    "path": "/compare/cerebrium-vs-lambda",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Lambda Cloud for AI agents, C 55.3 vs D 50",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
  },
  "tokens": {
    "markdown": 2100,
    "slim": 680
  },
  "version": 1
}
