{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 454,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T19:08:42.947721656Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 326,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 251,
          "p95ms24h": 326,
          "samples24h": 19,
          "samples30d": 19,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 19,
              "ok": 19
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:50:27.054213755Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T19:08:42.947721656Z"
      }
    },
    "answer": "Cerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics.",
    "b": {
      "slug": "koyeb",
      "name": "Koyeb",
      "vendor": "Koyeb",
      "vendorUrl": "https://www.koyeb.com",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Serverless platform for deploying applications and containers on CPU or GPU instances, with autoscaling and an API.",
      "url": "https://www.anchorterminal.com/tools/koyeb",
      "markdownUrl": "https://www.anchorterminal.com/tools/koyeb.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/koyeb.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/koyeb.json",
      "repo": "https://github.com/koyeb/koyeb-cli",
      "license": "Apache-2.0",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://app.koyeb.com/v1",
      "packages": [],
      "auth": "pat",
      "authNotes": "Personal access token created in the control panel (app.koyeb.com/user/settings/api) and sent as `Authorization: Bearer` to app.koyeb.com/v1. The CLI stores it from `koyeb login` and hides it in debug output. Organisation-scoped tokens can be created for teams.",
      "pricing": "freemium",
      "pricingNotes": "Pro is $29 a month with $10 of usage included, Scale $299 with $100, Enterprise custom. The Starter plan closed to new sign-ups after the Mistral AI agreement. Everything bills per second. GPUs an hour are RTX 4000 SFF Ada $0.50, L4 $0.70, RTX A6000 $0.75, L40S $1.20, A100 $1.60, A100 SXM $2.15, RTX PRO 6000 $2.20, H100 $2.50, H200 $3.00, 2x A100 $3.20, 2x H100 $5.00, 4x A100 $6.40, 8x H100 $20.00, 8x H200 $24.00. CPU services start at $0.000006 a second (https://www.koyeb.com/pricing, https://www.koyeb.com/blog/koyeb-is-joining-mistral-ai-to-build-the-future-of-ai-infrastructure).",
      "priceSummary": "$29 / mo",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 72,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://www.koyeb.com/docs",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "eu",
        "enterprise"
      ],
      "lastRelease": "2026-05-12",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 46.5,
        "grade": "D",
        "agentReady": false,
        "rank": 561,
        "ranked": true,
        "rankOf": 629,
        "categoryRank": 11,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 56,
          "maintenance": 23,
          "payments": 20,
          "reliability": 60,
          "schema": 41,
          "security": 50,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "H100 at $2.50 and H200 at $3.00 an hour, billed per second. Changelog silent since 27 February 2026 and no CLI release since 12 May 2026.",
        "bestFor": "Teams that want cheap H100 or H200 containers with scale to zero and an SLA, deployed from Git or an image without a vendor SDK.",
        "strengths": [
          "H100 at $2.50 and H200 at $3.00 an hour, billed per second",
          "Scale to zero after 5 idle minutes by default, configurable up to 12 hours on higher plans",
          "99.9 per cent uptime SLA published from the Pro plan",
          "List endpoints filter by name, type, status and region with `limit` and `offset`, and create and update take `dry_run`",
          "French company under French law with a published DPA"
        ],
        "weaknesses": [
          "Changelog silent since 27 February 2026 and no CLI release since 12 May 2026",
          "No free compute tier; Starter closed after the Mistral AI agreement",
          "No llms.txt, no downloadable OpenAPI file, no documented rate limits or error codes",
          "Audit logs and RBAC only on Enterprise, and tokens have no scopes",
          "Platform due to fold into Mistral Compute with no dated migration plan"
        ],
        "agentNotes": [
          "Pass `dry_run` on service create or update to validate the definition before anything deploys",
          "Page service lists with `limit` and `offset` and filter by `statuses` instead of fetching everything",
          "Expect the first request after deep sleep to take 1 to 5 seconds and retry once with a timeout",
          "Set `--min-scale 0` on GPU services so idle instances stop billing",
          "Re-check the Mistral transition before building anything long-lived on it"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 46.5
          }
        ],
        "editorialScores": {
          "ergonomics": 56,
          "maintenance": 23,
          "payments": 20,
          "reliability": 60,
          "schema": 41,
          "security": 50,
          "transparency": 50
        },
        "provenanceScore": 75
      },
      "connect": {
        "install": "brew install koyeb/tap/koyeb \u0026\u0026 koyeb login  # or go install github.com/koyeb/koyeb-cli/cmd/koyeb",
        "http": "curl \"https://app.koyeb.com/v1/services\" -H \"Authorization: Bearer $KOYEB_TOKEN\""
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/koyeb"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.5
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 3
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 1.6
        },
        {
          "item": "RTX PRO 6000",
          "unit": "gpu-hour",
          "usd": 2.2
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 1.2
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.7
        },
        {
          "item": "Pro plan",
          "unit": "month",
          "usd": 29,
          "note": "$10 of usage included"
        },
        {
          "item": "Scale plan",
          "unit": "month",
          "usd": 299,
          "note": "$100 of usage included"
        }
      ],
      "provenance": {
        "legalEntity": "Koyeb SAS",
        "domain": "koyeb.com",
        "domainRegistered": "2019-03-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.koyeb.com/docs/legal/terms",
        "privacy": "https://www.koyeb.com/docs/legal/data-processing-agreement",
        "statusPage": "https://status.koyeb.com",
        "changelog": "https://www.koyeb.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "The Master Services Agreement of 28 June 2024 names Koyeb, a simplified joint-stock company registered in Nanterre under 850 183 948, at 9 rue des Longs Prés, Boulogne-Billancourt, under French law.",
          "The site footer's privacy link points at the data processing agreement rather than a separate privacy policy.",
          "www.koyeb.com/.well-known/security.txt returns 404 and there's no llms.txt at www.koyeb.com or under /docs.",
          "Koyeb announced on 17 February 2026 that it is joining Mistral AI."
        ],
        "score": 75
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/koyeb.json",
      "live": {
        "slug": "koyeb",
        "probe": {
          "target": "https://app.koyeb.com/v1",
          "method": "get",
          "lastAt": "2026-10-08T19:08:50.781450165Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 169,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 72,
          "p95ms24h": 151,
          "samples24h": 272,
          "samples30d": 1933,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 217,
              "ok": 217
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.koyeb.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T17:50:48.615603505Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "koyeb/koyeb-cli",
            "version": "v5.12.0",
            "released": "2026-09-16",
            "seenAt": "2026-10-08T16:18:01.679616884Z"
          }
        ],
        "githubStars": 75,
        "securityTxt": {
          "url": "https://koyeb.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:33.165001961Z"
        },
        "domain": {
          "domain": "koyeb.com",
          "registered": "2019-03-11",
          "source": "https://rdap.verisign.com/com/v1/domain/koyeb.com",
          "checkedAt": "2026-10-04T13:09:21.713025238Z"
        },
        "pages": [
          {
            "url": "https://www.koyeb.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-08T18:28:38.583332732Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "7e4e89b36b63"
          },
          {
            "url": "https://www.koyeb.com/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-08T18:28:44.671933232Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "9d7cf517fa78"
          },
          {
            "url": "https://www.koyeb.com/docs/legal/data-processing-agreement",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-08T18:28:40.700270182Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6e08cc6e8919"
          },
          {
            "url": "https://www.koyeb.com/docs/legal/terms",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-08T18:28:42.728518755Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "089f524cdbd6"
          }
        ],
        "updatedAt": "2026-10-08T19:08:50.781450165Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "Model platform",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Koyeb",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "https://app.koyeb.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "Token",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "no",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "2026-05-12",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "no date given",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "no date given",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "72 stars",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Cerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics.",
        "question": "Which is better for AI agents, Cerebrium or Koyeb?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 70 against 41",
          "Security \u0026 auth, 60 against 50",
          "Payments \u0026 pricing, 30 against 20",
          "Maintenance \u0026 community, 75 against 23"
        ],
        "also": null,
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Reliability, 60 against 48",
          "Agent ergonomics, 56 against 49"
        ],
        "also": null,
        "goodFor": "Teams that want cheap H100 or H200 containers with scale to zero and an SLA, deployed from Git or an image without a vendor SDK.",
        "slug": "koyeb",
        "watchFor": "Changelog silent since 27 February 2026 and no CLI release since 12 May 2026"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-koyeb.json",
        "title": "Baseten vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-koyeb.json",
        "title": "Beam vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/beam-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-koyeb.json",
        "title": "CoreWeave vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-lambda.json",
        "title": "Koyeb vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-modal.json",
        "title": "Koyeb vs Modal",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-northflank.json",
        "title": "Koyeb vs Northflank",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.json",
        "title": "Koyeb vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-runpod.json",
        "title": "Koyeb vs Runpod",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-vast-ai.json",
        "title": "Koyeb vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 12,
        "cerebrium": 48,
        "edge": "koyeb",
        "key": "reliability",
        "koyeb": 60,
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 29,
        "cerebrium": 70,
        "edge": "cerebrium",
        "key": "schema",
        "koyeb": 41,
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 7,
        "cerebrium": 49,
        "edge": "koyeb",
        "key": "ergonomics",
        "koyeb": 56,
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 10,
        "cerebrium": 60,
        "edge": "cerebrium",
        "key": "security",
        "koyeb": 50,
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 10,
        "cerebrium": 30,
        "edge": "cerebrium",
        "key": "payments",
        "koyeb": 20,
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 52,
        "cerebrium": 75,
        "edge": "cerebrium",
        "key": "maintenance",
        "koyeb": 23,
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 0,
        "cerebrium": 63,
        "edge": "",
        "key": "transparency",
        "koyeb": 63,
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Cerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "koyeb": "H100 at $2.50 and H200 at $3.00 an hour, billed per second. Changelog silent since 27 February 2026 and no CLI release since 12 May 2026."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.min.md"
  },
  "markdown": "Cerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #454 of 629. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Koyeb: grade D, 46.5/100, rank #561 of 629. Markdown https://www.anchorterminal.com/tools/koyeb.md · JSON https://www.anchorterminal.com/api/v1/tools/koyeb.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAhead on:\n- Schema \u0026 documentation, 70 against 41\n- Security \u0026 auth, 60 against 50\n- Payments \u0026 pricing, 30 against 20\n- Maintenance \u0026 community, 75 against 23\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Koyeb (D)\n\nGood for: Teams that want cheap H100 or H200 containers with scale to zero and an SLA, deployed from Git or an image without a vendor SDK.\n\nAhead on:\n- Reliability, 60 against 48\n- Agent ergonomics, 56 against 49\n\nWatch for: Changelog silent since 27 February 2026 and no CLI release since 12 May 2026\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Koyeb | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 60 | Koyeb +12 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 41 | Cerebrium +29 |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 56 | Koyeb +7 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 50 | Cerebrium +10 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 20 | Cerebrium +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 23 | Cerebrium +52 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 63 | even |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **46.5 · D** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Koyeb |\n| --- | --- | --- |\n| Kind | Model platform | Model platform |\n| Vendor | Cerebrium Inc. | Koyeb |\n| Hosted endpoint | `https://rest.cerebrium.ai` | `https://app.koyeb.com/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | Token |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| Last release | 2026-09-16 | 2026-05-12 |\n| Terms last updated | no date given | no date given |\n| Privacy policy last updated | no date given | no date given |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | yes |\n| Terms restrict benchmarking | not found in the text | not found in the text |\n| Terms or service can change without notice | yes | yes |\n| Arbitration or class-action waiver | not found in the text | not found in the text |\n| Popularity | 920 PyPI/wk | 72 stars |\n| Agent reviews | none | 3.5/5 (2) |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Koyeb.** H100 at $2.50 and H200 at $3.00 an hour, billed per second. Changelog silent since 27 February 2026 and no CLI release since 12 May 2026.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Koyeb\n\n1. Pass `dry_run` on service create or update to validate the definition before anything deploys\n2. Page service lists with `limit` and `offset` and filter by `statuses` instead of fetching everything\n3. Expect the first request after deep sleep to take 1 to 5 seconds and retry once with a timeout\n4. Set `--min-scale 0` on GPU services so idle instances stop billing\n5. Re-check the Mistral transition before building anything long-lived on it\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Koyeb?\n\nCerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"koyeb\"}`. From a terminal: `anchor compare cerebrium koyeb`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/koyeb.json\n\n## Other comparisons with Cerebrium or Koyeb\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Koyeb](https://www.anchorterminal.com/compare/baseten-vs-koyeb.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Koyeb](https://www.anchorterminal.com/compare/beam-vs-koyeb.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n- [CoreWeave vs Koyeb](https://www.anchorterminal.com/compare/coreweave-vs-koyeb.md)\n- [Koyeb vs Lambda Cloud](https://www.anchorterminal.com/compare/koyeb-vs-lambda.md)\n- [Koyeb vs Modal](https://www.anchorterminal.com/compare/koyeb-vs-modal.md)\n- [Koyeb vs Northflank](https://www.anchorterminal.com/compare/koyeb-vs-northflank.md)\n- [Koyeb vs Replicate Deployments](https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.md)\n- [Koyeb vs Runpod](https://www.anchorterminal.com/compare/koyeb-vs-runpod.md)\n- [Koyeb vs Vast.ai](https://www.anchorterminal.com/compare/koyeb-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Koyeb",
        "url": ""
      }
    ],
    "description": "Cerebrium scores 55.3 (C) on agent readiness against Koyeb's 46.5 (D), and leads in 4 of 7 scored categories. Koyeb leads on reliability and agent ergonomics. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Koyeb D 46.5",
      "scores"
    ],
    "h1": "Cerebrium vs Koyeb",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-koyeb.png",
    "path": "/compare/cerebrium-vs-koyeb",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Koyeb for AI agents, C 55.3 vs D 46.5 | Anchor Terminal",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
  },
  "tokens": {
    "markdown": 2000,
    "slim": 530
  },
  "version": 1
}
