{
  "data": {
    "a": {
      "slug": "cerebrium",
      "name": "Cerebrium",
      "vendor": "Cerebrium Inc.",
      "vendorUrl": "https://www.cerebrium.ai",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Cerebrium is a serverless platform for running your own models and code on GPUs and CPUs. A CLI packages code into containers served as REST, streaming and WebSocket endpoints, managed through a REST API.",
      "url": "https://www.anchorterminal.com/tools/cerebrium",
      "markdownUrl": "https://www.anchorterminal.com/tools/cerebrium.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cerebrium.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cerebrium.json",
      "repo": "https://github.com/CerebriumAI/cerebrium",
      "license": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://rest.cerebrium.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cerebrium"
        }
      ],
      "auth": "api-key",
      "authNotes": "Two credentials. A service account token, created in the dashboard or over the API with an expiry of up to one year and a list of granted projects, authenticates the CLI (`CEREBRIUM_SERVICE_ACCOUNT_TOKEN`) and the management API at rest.cerebrium.ai as `Authorization: Bearer`. A project API key (a JWT) authenticates calls to deployed endpoints, and only when `cerebrium.toml` sets `disable_auth = false`. Signup and `cerebrium login` are browser flows.",
      "pricing": "freemium",
      "pricingNotes": "Hobby plan is $0 a month plus compute, Standard $100 a month plus compute, Enterprise on request. GPU, CPU and memory bill per second, from T4 at $0.000164 a second ($0.59 an hour) to H100 at $0.000944 ($3.40) and B200 at $0.00167 ($6.01). CPU $0.00000655 a vCPU-second, memory $0.00000222 a GB-second, storage $0.05 a GB-month after 100 GB free. Listed rates are for the default interruptible tier, and `protected` compute costs twice as much. Cold-start time is free, builds and model initialisation are billed. An account can start on Hobby without a contract. The pricing page does not say whether a card is needed or state a free compute allowance (https://www.cerebrium.ai/pricing, https://cerebrium.ai/docs/calculating-cost).",
      "priceSummary": "$0.0236 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs, the OpenAPI spec or the pricing page (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": 920,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://cerebrium.ai/docs",
      "llmsTxt": "https://cerebrium.ai/docs/llms.txt",
      "openapi": "https://s3.eu-west-1.amazonaws.com/www.cerebrium.ai/openapi_spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "serverless",
        "gpu",
        "cli",
        "openapi",
        "llms-txt",
        "python",
        "async-jobs",
        "multi-region",
        "status-page",
        "soc2",
        "hipaa"
      ],
      "lastRelease": "2026-09-16",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.3,
        "grade": "C",
        "agentReady": false,
        "rank": 512,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
        "bestFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "strengths": [
          "Per-second prices for ten GPU types published without a login, from T4 at $0.59 an hour to B200 at $6.01",
          "Public OpenAPI 3.0 spec for the management API at rest.cerebrium.ai, with 94 operations, plus llms.txt and Markdown docs",
          "Service account tokens carry an expiry of up to one year and a list of 1 to 50 granted projects",
          "Audit log of 22 actions with actor, IP address and outcome, readable over the API on Standard and Enterprise",
          "Status page with 11 components and 90 days of incident history, and five CLI releases between 7 August and 16 September 2026"
        ],
        "weaknesses": [
          "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it",
          "No request rate limits, 429 guidance, idempotency keys or SLA found in the reviewed documentation",
          "Several multi-hour degradations of the Inference API between 15 July and 2 September 2026, and a 10-minute outage on 21 July",
          "No deprecation policy, platform changelog or public subprocessor list found",
          "The CLI stores tokens in plaintext in `~/.cerebrium/config.yaml` with mode 0644, per its own SECURITY.md"
        ],
        "agentNotes": [
          "Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL",
          "Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser",
          "Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours",
          "Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller",
          "Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.3
          }
        ],
        "editorialScores": {
          "ergonomics": 49,
          "maintenance": 75,
          "payments": 30,
          "reliability": 48,
          "schema": 70,
          "security": 60,
          "transparency": 45
        },
        "provenanceScore": 80
      },
      "connect": {
        "install": "pip install cerebrium \u0026\u0026 cerebrium login",
        "http": "curl --location --request POST 'https://api.cerebrium.ai/v4/p-xxxxxxxx/{app-name}/{function}' \\\n  --header 'Authorization: Bearer \u003cJWT_TOKEN\u003e' \\\n  --header 'Content-Type: application/json' \\\n  --data '{\"function_param\": \"data\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/cerebrium"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.01,
          "note": "$0.00167 a second, interruptible tier"
        },
        {
          "item": "H200 141 GB",
          "unit": "gpu-hour",
          "usd": 4.2,
          "note": "$0.001166 a second, interruptible tier"
        },
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 3.4,
          "note": "$0.000944 a second, interruptible tier"
        },
        {
          "item": "RTX PRO 6000 96 GB",
          "unit": "gpu-hour",
          "usd": 2.5,
          "note": "$0.000694 a second, interruptible tier"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 2.1,
          "note": "$0.000583 a second, interruptible tier"
        },
        {
          "item": "A100 40 GB",
          "unit": "gpu-hour",
          "usd": 2,
          "note": "$0.000555 a second, interruptible tier"
        },
        {
          "item": "L40s 48 GB",
          "unit": "gpu-hour",
          "usd": 1.95,
          "note": "$0.000542 a second, interruptible tier"
        },
        {
          "item": "A10 24 GB",
          "unit": "gpu-hour",
          "usd": 1.1,
          "note": "$0.000306 a second, interruptible tier"
        },
        {
          "item": "L4 24 GB",
          "unit": "gpu-hour",
          "usd": 0.8,
          "note": "$0.000222 a second, interruptible tier"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.59,
          "note": "$0.000164 a second, interruptible tier"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.0236,
          "note": "$0.00000655 a vCPU-second, memory extra at $0.00000222 a GB-second"
        },
        {
          "item": "Persistent storage",
          "unit": "gb-month",
          "usd": 0.05,
          "note": "First 100 GB free"
        },
        {
          "item": "Standard plan",
          "unit": "month",
          "usd": 100,
          "note": "Plus compute"
        }
      ],
      "provenance": {
        "legalEntity": "Cerebrium Inc.",
        "domain": "cerebrium.ai",
        "domainRegistered": "2021-06-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cerebrium.ai/terms-of-service",
        "privacy": "https://www.cerebrium.ai/privacy",
        "statusPage": "https://status.cerebrium.ai",
        "changelog": "https://github.com/CerebriumAI/cerebrium/releases",
        "securityTxt": "none",
        "checked": "2026-10-08",
        "notes": [
          "The terms of service name Cerebrium Inc and say they are governed by the laws of the United Kingdom. The privacy policy names Cerebrium Inc. as data controller at 251 Little Falls Drive, Wilmington, Delaware.",
          "The terms of service are the only terms Cerebrium publishes. They cover accounts, subscriptions and the Service, and the OpenAPI spec names them as the API's licence. Neither document states a date.",
          "Deployed endpoints answer at api.cerebrium.ai and the management API at rest.cerebrium.ai. The OpenAPI file is served from an AWS S3 bucket.",
          "cerebrium.ai/.well-known/security.txt and cerebrium.ai/security.txt returned 404 on 8 October 2026. The CLI repository's SECURITY.md and the docs give security@cerebrium.ai.",
          "The changelog link is the CLI's GitHub releases. No platform changelog was found. Domain registration date from RDAP."
        ],
        "score": 80
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cerebrium.json",
      "live": {
        "slug": "cerebrium",
        "probe": {
          "target": "https://rest.cerebrium.ai",
          "method": "get",
          "lastAt": "2026-10-08T20:21:09.854161324Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 274,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 253,
          "p95ms24h": 326,
          "samples24h": 32,
          "samples30d": 32,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cerebrium.ai",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:38:18.718586754Z"
        },
        "pages": [
          {
            "url": "https://www.cerebrium.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:57.20274404Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "be1ada48db42"
          },
          {
            "url": "https://www.cerebrium.ai/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:26:59.281048683Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "017a8d228768"
          },
          {
            "url": "https://www.cerebrium.ai/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:27:01.554163635Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a99804810a11"
          }
        ],
        "updatedAt": "2026-10-08T20:21:09.854161324Z"
      }
    },
    "answer": "Replicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community.",
    "b": {
      "slug": "replicate-deploy",
      "name": "Replicate Deployments",
      "vendor": "Replicate",
      "vendorUrl": "https://replicate.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Replicate's service for deploying and running custom models.",
      "url": "https://www.anchorterminal.com/tools/replicate-deploy",
      "markdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json",
      "repo": "https://github.com/replicate/cog",
      "license": "Apache-2.0",
      "transports": [
        "http",
        "sse",
        "stdio"
      ],
      "remoteUrl": "https://api.replicate.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "replicate"
        },
        {
          "registry": "pypi",
          "name": "replicate"
        },
        {
          "registry": "npm",
          "name": "replicate-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API token on every call to api.replicate.com. `cog push` uses the same token to upload a model image. The hosted MCP at https://mcp.replicate.com/sse asks for the token in a browser flow and holds it for the client; the local `replicate-mcp` package reads `REPLICATE_API_TOKEN`.",
      "pricing": "usage",
      "pricingNotes": "Private models and deployments bill per second for the whole time an instance is up, set-up and idle included, from prepaid credit or monthly in arrears. CPU $0.000100 a second ($0.36 an hour), T4 $0.000225 ($0.81), L40S $0.000975 ($3.51), A100 80 GB $0.001400 ($5.04), H100 $0.001525 ($5.49), 2x L40S $0.001950 ($7.02), 2x A100 $0.002800 ($10.08). 2x H100 ($10.98), 4x and 8x L40S, A100 and H100 up to $43.92 an hour need a committed-spend contract. Fast-booting fine-tunes bill only while active. Public models bill only active time and not failures (https://replicate.com/pricing, https://replicate.com/docs/topics/billing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 9500,
        "npmWeekly": 634116,
        "pypiWeekly": 386704,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://replicate.com/docs/topics/deployments",
      "llmsTxt": "https://replicate.com/docs/llms.txt",
      "openapi": "https://api.replicate.com/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "mcp",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "async-jobs",
        "webhooks",
        "open-source"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63.6,
        "grade": "B",
        "agentReady": false,
        "rank": 303,
        "ranked": true,
        "rankOf": 722,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 78
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.",
        "bestFor": "Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.",
        "strengths": [
          "OpenAPI file, llms.txt and an MCP server with a two-tool code mode",
          "Deployment min and max instances settable over the API, 0 allowed",
          "API prediction data deleted after one hour by default",
          "Published limits, 600 prediction creates and 3,000 other calls a minute",
          "Leaked tokens found on GitHub are disabled automatically"
        ],
        "weaknesses": [
          "Private instances bill set-up and idle time, H100 at $5.49 an hour",
          "API tokens have no scopes, expiry or audit log",
          "Changelog silent since 21 April 2026",
          "Only T4, L40S, A100 and H100, and more than 2 GPUs needs a committed-spend contract",
          "Two September 2026 incidents ran 15 and 20 hours, both marked minor"
        ],
        "agentNotes": [
          "List `GET /v1/hardware` first and use the returned `sku` in the deployment body",
          "Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not",
          "Send `Prefer: wait` on deployment predictions to block instead of polling",
          "Copy outputs within an hour; API prediction data is deleted after that",
          "Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63.6
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 69
        },
        "provenanceScore": 87
      },
      "connect": {
        "install": "pip install cog replicate",
        "http": "curl -X POST \"https://api.replicate.com/v1/deployments/$REPLICATE_OWNER/my-deployment/predictions\" \\\n  -H \"Authorization: Bearer $REPLICATE_API_TOKEN\" -H \"Content-Type: application/json\" -H \"Prefer: wait\" \\\n  -d '{\"input\":{\"prompt\":\"hello\"}}'",
        "claudeCode": "claude mcp add replicate https://mcp.replicate.com/sse --transport sse --scope user",
        "config": {
          "mcpServers": {
            "replicate": {
              "args": [
                "-y",
                "replicate-mcp"
              ],
              "command": "npx",
              "env": {
                "REPLICATE_API_TOKEN": "${REPLICATE_API_TOKEN}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/replicate-deploy"
      },
      "sameCompany": [
        "replicate-image",
        "replicate-musicgen"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.49,
          "note": "$0.001525 a second, including set-up and idle"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.04,
          "note": "$0.001400 a second"
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 3.51,
          "note": "$0.000975 a second"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.81,
          "note": "$0.000225 a second"
        }
      ],
      "provenance": {
        "legalEntity": "Replicate, LLC",
        "domain": "replicate.com",
        "domainRegistered": "1998-05-26",
        "domainNote": "replicate.com was registered in 1998, long before Replicate the company existed.",
        "endpointOnVendorDomain": true,
        "terms": "https://replicate.com/terms",
        "privacy": "https://replicate.com/privacy",
        "statusPage": "https://replicatestatus.com",
        "changelog": "https://replicate.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms last updated 2026-04-01 name Replicate, LLC as the contracting party.",
          "replicatestatus.com redirects to Cloudflare's status page filtered to Replicate.",
          "Replicate's hosted image and music models are listed separately under image generation and music generation."
        ],
        "score": 87
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/replicate-deploy.json",
      "live": {
        "slug": "replicate-deploy",
        "probe": {
          "target": "https://api.replicate.com/v1",
          "method": "get",
          "lastAt": "2026-10-08T20:21:25.346392865Z",
          "lastOk": true,
          "lastStatus": 401,
          "lastMs": 283,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 143,
          "p95ms24h": 328,
          "samples24h": 272,
          "samples30d": 1946,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 230,
              "ok": 230
            }
          ]
        },
        "vendorStatus": {
          "page": "https://replicatestatus.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-08T19:39:07.002266541Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "replicate/cog",
            "version": "v0.23.0",
            "released": "2026-09-22",
            "seenAt": "2026-10-08T16:27:09.246140372Z"
          },
          {
            "registry": "npm",
            "name": "replicate",
            "version": "1.4.0",
            "seenAt": "2026-10-08T16:27:06.602404036Z"
          },
          {
            "registry": "npm",
            "name": "replicate-mcp",
            "version": "0.9.0",
            "seenAt": "2026-10-08T16:27:09.031273281Z"
          },
          {
            "registry": "pypi",
            "name": "replicate",
            "version": "1.0.7",
            "released": "2025-05-27",
            "seenAt": "2026-10-08T16:27:07.620263226Z"
          }
        ],
        "githubStars": 9487,
        "npmWeekly": 704939,
        "pypiWeekly": 356876,
        "securityTxt": {
          "url": "https://replicate.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:39:00.17768337Z"
        },
        "llmsTxt": {
          "url": "https://replicate.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:48.463855942Z"
        },
        "domain": {
          "domain": "replicate.com",
          "registered": "1998-05-26",
          "source": "https://rdap.verisign.com/com/v1/domain/replicate.com",
          "checkedAt": "2026-10-04T13:07:04.742407865Z"
        },
        "pages": [
          {
            "url": "https://replicate.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-08T18:23:40.627013731Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "490f4836aca3"
          },
          {
            "url": "https://replicate.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:43.001722065Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3f1305f154be"
          },
          {
            "url": "https://replicate.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:45.469231625Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "8e299fbc64eb"
          },
          {
            "url": "https://replicate.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:46.848375869Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ea48efe3382b"
          }
        ],
        "updatedAt": "2026-10-08T20:21:25.346392865Z"
      }
    },
    "facts": [
      {
        "a": "Model platform",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Cerebrium Inc.",
        "b": "Replicate",
        "name": "Vendor"
      },
      {
        "a": "https://rest.cerebrium.ai",
        "b": "https://api.replicate.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, SSE (legacy), stdio",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cerebrium's terms of service. The CLI is MIT",
        "b": "Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-16",
        "b": "2026-09-22",
        "name": "Last release"
      },
      {
        "a": "no date given",
        "b": "2026-04-01",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2026-04-01",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "920 PyPI/wk",
        "b": "9.5k stars, 634k npm/wk, 387k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Replicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community.",
        "question": "Which is better for AI agents, Cerebrium or Replicate Deployments?"
      },
      {
        "answer": "Yes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Replicate Deployments at https://api.replicate.com/v1.",
        "question": "Can an agent call Cerebrium and Replicate Deployments without installing anything?"
      },
      {
        "answer": "No open-source release is listed for Cerebrium. Replicate Deployments is open source (Apache-2.0).",
        "question": "Are Cerebrium and Replicate Deployments open source?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Security \u0026 auth, 60 against 40",
          "Maintenance \u0026 community, 75 against 70"
        ],
        "also": null,
        "goodFor": "Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.",
        "slug": "cerebrium",
        "watchFor": "`disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it"
      },
      {
        "aheadOn": [
          "Reliability, 75 against 48",
          "Schema \u0026 documentation, 85 against 70",
          "Agent ergonomics, 68 against 49",
          "Transparency \u0026 trust, 78 against 63"
        ],
        "also": [
          "Runs on your own machine",
          "Open source"
        ],
        "goodFor": "Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.",
        "slug": "replicate-deploy",
        "watchFor": "Private instances bill set-up and idle time, H100 at $5.49 an hour"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium.json",
        "title": "Baseten vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.json",
        "title": "Baseten vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-cerebrium.json",
        "title": "Beam vs Cerebrium",
        "url": "https://www.anchorterminal.com/compare/beam-vs-cerebrium"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.json",
        "title": "Beam vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.json",
        "title": "Cerebrium vs CoreWeave",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-coreweave"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.json",
        "title": "Cerebrium vs Koyeb",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-koyeb"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda.json",
        "title": "Cerebrium vs Lambda Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-lambda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-modal.json",
        "title": "Cerebrium vs Modal",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-modal"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank.json",
        "title": "Cerebrium vs Northflank",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod.json",
        "title": "Cerebrium vs Runpod",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.json",
        "title": "Cerebrium vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy.json",
        "title": "CoreWeave vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.json",
        "title": "Koyeb vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.json",
        "title": "Lambda Cloud vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.json",
        "title": "Modal vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.json",
        "title": "Northflank vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.json",
        "title": "Replicate Deployments vs Runpod",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.json",
        "title": "Replicate Deployments vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai"
      }
    ],
    "scores": [
      {
        "by": 27,
        "cerebrium": 48,
        "edge": "replicate-deploy",
        "key": "reliability",
        "name": "Reliability",
        "replicate-deploy": 75,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 15,
        "cerebrium": 70,
        "edge": "replicate-deploy",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "replicate-deploy": 85,
        "weight": 13
      },
      {
        "by": 19,
        "cerebrium": 49,
        "edge": "replicate-deploy",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "replicate-deploy": 68,
        "weight": 13
      },
      {
        "by": 20,
        "cerebrium": 60,
        "edge": "cerebrium",
        "key": "security",
        "name": "Security \u0026 auth",
        "replicate-deploy": 40,
        "weight": 14
      },
      {
        "by": 0,
        "cerebrium": 30,
        "edge": "",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "replicate-deploy": 30,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 5,
        "cerebrium": 75,
        "edge": "cerebrium",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "replicate-deploy": 70,
        "weight": 7
      },
      {
        "by": 15,
        "cerebrium": 63,
        "edge": "replicate-deploy",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "replicate-deploy": 78,
        "weight": 7
      }
    ],
    "summary": "Replicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community. Both do compute gpu.",
    "verdicts": {
      "cerebrium": "Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.",
      "replicate-deploy": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy",
    "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md",
    "slim": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.min.md"
  },
  "markdown": "Replicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community. Both do compute gpu.\n\n- Cerebrium: grade C, 55.3/100, rank #512 of 722. Markdown https://www.anchorterminal.com/tools/cerebrium.md · JSON https://www.anchorterminal.com/api/v1/tools/cerebrium.json\n- Replicate Deployments: grade B, 63.6/100, rank #303 of 722. Markdown https://www.anchorterminal.com/tools/replicate-deploy.md · JSON https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Which one, for what\n\n### Cerebrium (C)\n\nGood for: Teams serving their own models as real-time endpoints (voice, LLM, image) who want per-second billing, multi-region placement and a scriptable management API.\n\nAhead on:\n- Security \u0026 auth, 60 against 40\n- Maintenance \u0026 community, 75 against 70\n\nWatch for: `disable_auth` defaults to true, so a deployed endpoint answers without a token unless the owner changes it\n\n### Replicate Deployments (B)\n\nGood for: Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.\n\nAhead on:\n- Reliability, 75 against 48\n- Schema \u0026 documentation, 85 against 70\n- Agent ergonomics, 68 against 49\n- Transparency \u0026 trust, 78 against 63\n\nAlso in its favour:\n- Runs on your own machine\n- Open source\n\nWatch for: Private instances bill set-up and idle time, H100 at $5.49 an hour\n\n\n## Score by category\n\n| Category | Weight | Cerebrium | Replicate Deployments | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 48 | 75 | Replicate Deployments +27 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 70 | 85 | Replicate Deployments +15 |\n| Agent ergonomics | 13% (16.2 this run) | 49 | 68 | Replicate Deployments +19 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 40 | Cerebrium +20 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 30 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 70 | Cerebrium +5 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 63 | 78 | Replicate Deployments +15 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **55.3 · C** | **63.6 · B** | |\n\n## Facts side by side\n\n| Fact | Cerebrium | Replicate Deployments |\n| --- | --- | --- |\n| Kind | Model platform | HTTP API |\n| Vendor | Cerebrium Inc. | Replicate |\n| Hosted endpoint | `https://rest.cerebrium.ai` | `https://api.replicate.com/v1` |\n| Transports | HTTP | HTTP, SSE (legacy), stdio |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary service under Cerebrium's terms of service. The CLI is MIT | Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-16 | 2026-09-22 |\n| Terms last updated | no date given | 2026-04-01 |\n| Privacy policy last updated | no date given | 2026-04-01 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | not found in the text |\n| Terms restrict benchmarking | not found in the text | not found in the text |\n| Terms or service can change without notice | yes | yes |\n| Arbitration or class-action waiver | not found in the text | yes |\n| Popularity | 920 PyPI/wk | 9.5k stars, 634k npm/wk, 387k PyPI/wk |\n| Agent reviews | none | 3/5 (2) |\n\n## Verdicts\n\n**Cerebrium.** Per-second GPU prices are public, a 94-operation OpenAPI spec covers the management API, and service account tokens expire and are limited to named projects. Deployed endpoints are callable without a token unless `disable_auth = false` is set, and no request rate limits, 429 handling or SLA were found in the reviewed documentation.\n\n**Replicate Deployments.** OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.\n\n## Before you call either\n\n### Cerebrium\n\n1. Set `disable_auth = false` in `cerebrium.toml` before deploying. The default leaves the endpoint callable by anyone with the URL\n2. Authenticate headless with `CEREBRIUM_SERVICE_ACCOUNT_TOKEN`. `cerebrium login` opens a browser\n3. Raise `response_grace_period` for long work. It defaults to 15 minutes and async runs stop at 12 hours\n4. Send `?async=true` to get a `run_id` with HTTP 202, and add `webhookEndpoint` because async calls return no result to the caller\n5. Check the plan before choosing hardware. A100, H100, H200, B200 and RTX PRO 6000 need Standard, and `protected` compute bills at twice the listed rate\n\n### Replicate Deployments\n\n1. List `GET /v1/hardware` first and use the returned `sku` in the deployment body\n2. Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not\n3. Send `Prefer: wait` on deployment predictions to block instead of polling\n4. Copy outputs within an hour; API prediction data is deleted after that\n5. Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute\n\n## Questions\n\n### Which is better for AI agents, Cerebrium or Replicate Deployments?\n\nReplicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community.\n\n### Can an agent call Cerebrium and Replicate Deployments without installing anything?\n\nYes. Cerebrium has a hosted endpoint at https://rest.cerebrium.ai and Replicate Deployments at https://api.replicate.com/v1.\n\n### Are Cerebrium and Replicate Deployments open source?\n\nNo open-source release is listed for Cerebrium. Replicate Deployments is open source (Apache-2.0).\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cerebrium\", \"b\": \"replicate-deploy\"}`. From a terminal: `anchor compare cerebrium replicate-deploy`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cerebrium.json and https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Other comparisons with Cerebrium or Replicate Deployments\n\n- [Baseten vs Cerebrium](https://www.anchorterminal.com/compare/baseten-vs-cerebrium.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Beam vs Cerebrium](https://www.anchorterminal.com/compare/beam-vs-cerebrium.md)\n- [Beam vs Replicate Deployments](https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.md)\n- [Cerebrium vs CoreWeave](https://www.anchorterminal.com/compare/cerebrium-vs-coreweave.md)\n- [Cerebrium vs Koyeb](https://www.anchorterminal.com/compare/cerebrium-vs-koyeb.md)\n- [Cerebrium vs Lambda Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-lambda.md)\n- [Cerebrium vs Modal](https://www.anchorterminal.com/compare/cerebrium-vs-modal.md)\n- [Cerebrium vs Northflank](https://www.anchorterminal.com/compare/cerebrium-vs-northflank.md)\n- [Cerebrium vs Runpod](https://www.anchorterminal.com/compare/cerebrium-vs-runpod.md)\n- [Cerebrium vs Vast.ai](https://www.anchorterminal.com/compare/cerebrium-vs-vast-ai.md)\n- [CoreWeave vs Replicate Deployments](https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy.md)\n- [Koyeb vs Replicate Deployments](https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.md)\n- [Lambda Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md)\n- [Modal vs Replicate Deployments](https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.md)\n- [Northflank vs Replicate Deployments](https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.md)\n- [Replicate Deployments vs Runpod](https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.md)\n- [Replicate Deployments vs Vast.ai](https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-08",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cerebrium vs Replicate Deployments",
        "url": ""
      }
    ],
    "description": "Replicate Deployments scores 63.6 (B) on agent readiness against Cerebrium's 55.3 (C), and leads in 4 of 7 scored categories. Cerebrium leads on security \u0026 auth and maintenance \u0026 community. Both do compute gpu. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Cerebrium C 55.3",
      "Replicate Deployments B 63.6",
      "scores"
    ],
    "h1": "Cerebrium vs Replicate Deployments",
    "image": "https://www.anchorterminal.com/assets/og/compare-cerebrium-vs-replicate-deploy.png",
    "path": "/compare/cerebrium-vs-replicate-deploy",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cerebrium vs Replicate Deployments for AI agents, C 55.3 vs B 63.6",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
  },
  "tokens": {
    "markdown": 2250,
    "slim": 680
  },
  "version": 1
}
