{
  "data": {
    "a": {
      "slug": "beam",
      "name": "Beam",
      "vendor": "Beam",
      "vendorUrl": "https://www.beam.cloud",
      "kind": "platform",
      "category": "gpu-compute",
      "summary": "Serverless GPU endpoints, task queues, functions, pods and sandboxes from Python decorators, on the open-source beta9 runtime.",
      "url": "https://www.anchorterminal.com/tools/beam",
      "markdownUrl": "https://www.anchorterminal.com/tools/beam.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/beam.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/beam.json",
      "repo": "https://github.com/beam-cloud/beta9",
      "license": "AGPL-3.0",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://app.beam.cloud/api/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "beam-client"
        }
      ],
      "auth": "api-key",
      "authNotes": "API token from platform.beam.cloud, read from `BEAM_TOKEN` by the SDK and CLI or stored by `beam login` in `~/.beam/config.ini`, with named contexts for several workspaces. Deployed endpoints take the same token as `Authorization: Bearer`. The TypeScript SDK sets `beamOpts.token` server-side.",
      "pricing": "freemium",
      "pricingNotes": "Developer plan is free plus usage, Team $89 a month plus usage, Growth on request. Billed by the millisecond only while a container runs, which includes `on_start` and `keep_warm_seconds`; cold starts and image pulls aren't billed. Serverless GPUs are RTX 4090 24 GB $0.000192 a second ($0.69 an hour), RTX 5090 32 GB $0.000303 ($1.09), H100 PCIe 80 GB $0.000972 ($3.50). Reserved on-demand machines from $0.44 an hour (RTX 4090), $1.36 (A100 80 GB), $1.83 (H100 PCIe), $2.09 (H200), $4.11 (B200), billed until released even when idle. CPU $0.0000125 a core-second and RAM $0.0000021 a GiB-second on CPU-only work, $0.000105 and $0.0000055 when attached to a GPU, $0.0000375 and $0.0000064 in sandboxes. Storage 1 TB included, then $0.021 a GB-month (https://www.beam.cloud/pricing, https://docs.beam.cloud/v2/resources/pricing-and-billing).",
      "priceSummary": "$0.045 / vCPU-hr",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 1800,
        "npmWeekly": null,
        "pypiWeekly": 8481,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.beam.cloud",
      "llmsTxt": "https://docs.beam.cloud/llms.txt",
      "capabilities": [
        "compute.gpu",
        "compute.serverless",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "open-source",
        "self-hosted",
        "python",
        "typescript",
        "llms-txt",
        "async-jobs"
      ],
      "lastRelease": "2026-10-01",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.5,
        "grade": "C",
        "agentReady": false,
        "rank": 313,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 5,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 58,
          "maintenance": 85,
          "payments": 40,
          "reliability": 55,
          "schema": 58,
          "security": 50,
          "transparency": 74
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -2,
        "negativeNotes": [
          "-2: `beam deploy --format json` writes the full workspace bearer token into the JSON logs array that CI systems keep. Filed by a Beam engineer on 2026-08-04 and still open on 2026-10-01 (https://github.com/beam-cloud/beta9/issues/1828)"
        ],
        "verdict": "Per-millisecond billing with cold starts and image pulls free, H100 PCIe at $3.50 and RTX 4090 at $0.69 an hour. No published request rate limits, 429 handling or SLA.",
        "strengths": [
          "Per-millisecond billing with cold starts and image pulls free, H100 PCIe at $3.50 and RTX 4090 at $0.69 an hour",
          "Free Developer plan with no card required",
          "beta9, the engine the hosted cloud runs on, is AGPL-3.0 and self-hostable",
          "Workspace REST API and an official MCP server, local or remote, on the same token",
          "Privacy policy with retention periods per category and a published subprocessor list"
        ],
        "weaknesses": [
          "No published request rate limits, 429 handling or SLA",
          "No public changelog or deprecation notices; releases show up only on PyPI and GitHub",
          "Tokens have no documented scopes or read-only mode",
          "`beam deploy --format json` leaks the workspace token into its logs, open since 4 August 2026",
          "Status page silent since June 2025 despite sign-up failures reported on GitHub in August 2026"
        ],
        "agentNotes": [
          "Check the response body for `ok: false` on gateway calls; a failure can arrive as HTTP 200",
          "Don't pipe `beam deploy --format json` output into CI logs, since it contains the workspace token",
          "Route anything over 180 seconds to a task queue and poll the task instead of holding the endpoint request",
          "Set `keep_warm_seconds` deliberately; the 180-second endpoint default bills three minutes of GPU after every call",
          "Pass `gpu=[\"RTX4090\", \"A10G\"]` so a job still schedules when one type is out"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.5
          }
        ],
        "editorialScores": {
          "ergonomics": 58,
          "maintenance": 85,
          "payments": 40,
          "reliability": 55,
          "schema": 58,
          "security": 50,
          "transparency": 71
        },
        "provenanceScore": 76
      },
      "connect": {
        "install": "pip install beam-client \u0026\u0026 beam login",
        "http": "curl -X POST \"https://my-function-$BEAM_DEPLOYMENT_ID-v1.app.beam.cloud\" \\\n  -H \"Authorization: Bearer $BEAM_TOKEN\" -H \"Content-Type: application/json\" \\\n  -d '{\"x\":10}'"
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/beam"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 PCIe 80 GB serverless",
          "unit": "gpu-hour",
          "usd": 3.5,
          "note": "$0.000972 a second"
        },
        {
          "item": "RTX 5090 32 GB serverless",
          "unit": "gpu-hour",
          "usd": 1.09,
          "note": "$0.000303 a second"
        },
        {
          "item": "RTX 4090 24 GB serverless",
          "unit": "gpu-hour",
          "usd": 0.69,
          "note": "$0.000192 a second"
        },
        {
          "item": "B200 180 GB reserved machine",
          "unit": "gpu-hour",
          "usd": 4.11,
          "note": "From price, billed while reserved"
        },
        {
          "item": "H200 141 GB reserved machine",
          "unit": "gpu-hour",
          "usd": 2.09,
          "note": "From price, billed while reserved"
        },
        {
          "item": "A100 80 GB reserved machine",
          "unit": "gpu-hour",
          "usd": 1.36,
          "note": "From price, billed while reserved"
        },
        {
          "item": "CPU-only compute",
          "unit": "vcpu-hour",
          "usd": 0.045,
          "note": "$0.0000125 a core-second, RAM extra at $0.0000021 a GiB-second"
        },
        {
          "item": "Team plan",
          "unit": "month",
          "usd": 89,
          "note": "Plus usage"
        }
      ],
      "provenance": {
        "legalEntity": "Smartshare, Inc.",
        "domain": "beam.cloud",
        "domainRegistered": "2019-07-31",
        "endpointOnVendorDomain": true,
        "terms": "https://docs.beam.cloud/v2/security/terms-and-conditions",
        "privacy": "https://docs.beam.cloud/v2/security/privacy-policy",
        "statusPage": "https://status.beam.cloud",
        "changelog": "",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms dated 14 September 2026 name Smartshare, Inc., a Delaware corporation doing business as Beam. The site footer reads © 2026 Smartshare, Inc.",
          "Deployed endpoints run on app.beam.cloud, a subdomain of the vendor domain. There's no public REST base URL for deploying.",
          "www.beam.cloud/.well-known/security.txt and www.beam.cloud/terms return 404; the legal pages live under docs.beam.cloud.",
          "No public changelog found; releases show up as commits in the beam-client and beta9 repositories."
        ],
        "score": 76
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/beam.json",
      "live": {
        "slug": "beam",
        "probe": {
          "target": "https://app.beam.cloud/api/v1",
          "method": "get",
          "lastAt": "2026-10-04T22:35:19.53069957Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 265,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 259,
          "p95ms24h": 322,
          "samples24h": 272,
          "samples30d": 627,
          "days": [
            {
              "date": "2026-10-02",
              "probes": 100,
              "ok": 100
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 256,
              "ok": 256
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.beam.cloud",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T22:33:47.090252109Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "beam-cloud/beta9",
            "version": "worker-0.1.781",
            "released": "2026-10-03",
            "seenAt": "2026-10-04T16:22:06.02863332Z"
          },
          {
            "registry": "pypi",
            "name": "beam-client",
            "version": "0.2.217",
            "released": "2026-10-02",
            "seenAt": "2026-10-04T16:22:05.839399711Z"
          }
        ],
        "githubStars": 1802,
        "pypiWeekly": 9877,
        "securityTxt": {
          "url": "https://beam.cloud/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:16:04.647716161Z"
        },
        "llmsTxt": {
          "url": "https://docs.beam.cloud/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:19.657901017Z"
        },
        "domain": {
          "domain": "beam.cloud",
          "registered": "2019-07-31",
          "source": "https://rdap.registry.cloud/rdap/domain/beam.cloud",
          "checkedAt": "2026-10-04T13:04:15.987834596Z"
        },
        "pages": [
          {
            "url": "https://docs.beam.cloud/v2/resources/pricing-and-billing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:14.08690057Z",
            "changedAt": "2026-10-04T15:43:14.08690057Z",
            "fingerprint": "9430ae529da0"
          },
          {
            "url": "https://www.beam.cloud/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-04T15:49:25.671086961Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "dedbd2b07b0a"
          },
          {
            "url": "https://docs.beam.cloud/v2/security/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:16.313382922Z",
            "changedAt": "2026-10-04T15:43:16.313382922Z",
            "fingerprint": "4072b33654b6"
          },
          {
            "url": "https://docs.beam.cloud/v2/security/terms-and-conditions",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:18.294864193Z",
            "changedAt": "2026-10-04T15:43:18.294864193Z",
            "fingerprint": "988292e328cf"
          }
        ],
        "updatedAt": "2026-10-04T22:35:19.53069957Z"
      }
    },
    "b": {
      "slug": "replicate-deploy",
      "name": "Replicate Deployments",
      "vendor": "Replicate",
      "vendorUrl": "https://replicate.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Replicate's service for deploying and running custom models.",
      "url": "https://www.anchorterminal.com/tools/replicate-deploy",
      "markdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json",
      "repo": "https://github.com/replicate/cog",
      "license": "Apache-2.0",
      "transports": [
        "http",
        "sse",
        "stdio"
      ],
      "remoteUrl": "https://api.replicate.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "replicate"
        },
        {
          "registry": "pypi",
          "name": "replicate"
        },
        {
          "registry": "npm",
          "name": "replicate-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API token on every call to api.replicate.com. `cog push` uses the same token to upload a model image. The hosted MCP at https://mcp.replicate.com/sse asks for the token in a browser flow and holds it for the client; the local `replicate-mcp` package reads `REPLICATE_API_TOKEN`.",
      "pricing": "usage",
      "pricingNotes": "Private models and deployments bill per second for the whole time an instance is up, set-up and idle included, from prepaid credit or monthly in arrears. CPU $0.000100 a second ($0.36 an hour), T4 $0.000225 ($0.81), L40S $0.000975 ($3.51), A100 80 GB $0.001400 ($5.04), H100 $0.001525 ($5.49), 2x L40S $0.001950 ($7.02), 2x A100 $0.002800 ($10.08). 2x H100 ($10.98), 4x and 8x L40S, A100 and H100 up to $43.92 an hour need a committed-spend contract. Fast-booting fine-tunes bill only while active. Public models bill only active time and not failures (https://replicate.com/pricing, https://replicate.com/docs/topics/billing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 9500,
        "npmWeekly": 634116,
        "pypiWeekly": 386704,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://replicate.com/docs/topics/deployments",
      "llmsTxt": "https://replicate.com/docs/llms.txt",
      "openapi": "https://api.replicate.com/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "mcp",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "async-jobs",
        "webhooks",
        "open-source"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63.7,
        "grade": "B",
        "agentReady": false,
        "rank": 197,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 3,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 80
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.",
        "strengths": [
          "OpenAPI file, llms.txt and an MCP server with a two-tool code mode",
          "Deployment min and max instances settable over the API, 0 allowed",
          "API prediction data deleted after one hour by default",
          "Published limits, 600 prediction creates and 3,000 other calls a minute",
          "Leaked tokens found on GitHub are disabled automatically"
        ],
        "weaknesses": [
          "Private instances bill set-up and idle time, H100 at $5.49 an hour",
          "API tokens have no scopes, expiry or audit log",
          "Changelog silent since 21 April 2026",
          "Only T4, L40S, A100 and H100, and more than 2 GPUs needs a committed-spend contract",
          "Two September 2026 incidents ran 15 and 20 hours, both marked minor"
        ],
        "agentNotes": [
          "List `GET /v1/hardware` first and use the returned `sku` in the deployment body",
          "Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not",
          "Send `Prefer: wait` on deployment predictions to block instead of polling",
          "Copy outputs within an hour; API prediction data is deleted after that",
          "Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63.7
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 69
        },
        "provenanceScore": 90
      },
      "connect": {
        "install": "pip install cog replicate",
        "http": "curl -X POST \"https://api.replicate.com/v1/deployments/$REPLICATE_OWNER/my-deployment/predictions\" \\\n  -H \"Authorization: Bearer $REPLICATE_API_TOKEN\" -H \"Content-Type: application/json\" -H \"Prefer: wait\" \\\n  -d '{\"input\":{\"prompt\":\"hello\"}}'",
        "claudeCode": "claude mcp add replicate https://mcp.replicate.com/sse --transport sse --scope user",
        "config": {
          "mcpServers": {
            "replicate": {
              "args": [
                "-y",
                "replicate-mcp"
              ],
              "command": "npx",
              "env": {
                "REPLICATE_API_TOKEN": "${REPLICATE_API_TOKEN}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/replicate-deploy"
      },
      "sameCompany": [
        "replicate-image",
        "replicate-musicgen"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.49,
          "note": "$0.001525 a second, including set-up and idle"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.04,
          "note": "$0.001400 a second"
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 3.51,
          "note": "$0.000975 a second"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.81,
          "note": "$0.000225 a second"
        }
      ],
      "provenance": {
        "legalEntity": "Replicate, LLC",
        "domain": "replicate.com",
        "domainRegistered": "1998-05-26",
        "domainNote": "replicate.com was registered in 1998, long before Replicate the company existed.",
        "endpointOnVendorDomain": true,
        "terms": "https://replicate.com/terms",
        "privacy": "https://replicate.com/privacy",
        "statusPage": "https://replicatestatus.com",
        "changelog": "https://replicate.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms last updated 2026-04-01 name Replicate, LLC as the contracting party.",
          "replicatestatus.com redirects to Cloudflare's status page filtered to Replicate.",
          "Replicate's hosted image and music models are listed separately under image generation and music generation."
        ],
        "score": 90
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/replicate-deploy.json",
      "live": {
        "slug": "replicate-deploy",
        "probe": {
          "target": "https://api.replicate.com/v1",
          "method": "get",
          "lastAt": "2026-10-04T22:35:30.021478789Z",
          "lastOk": true,
          "lastStatus": 401,
          "lastMs": 123,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 138,
          "p95ms24h": 328,
          "samples24h": 272,
          "samples30d": 884,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 256,
              "ok": 256
            }
          ]
        },
        "vendorStatus": {
          "page": "https://replicatestatus.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-04T21:40:25.933184888Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "replicate/cog",
            "version": "v0.23.0",
            "released": "2026-09-22",
            "seenAt": "2026-10-04T16:38:03.363821386Z"
          },
          {
            "registry": "npm",
            "name": "replicate",
            "version": "1.4.0",
            "seenAt": "2026-10-04T16:38:01.019271447Z"
          },
          {
            "registry": "npm",
            "name": "replicate-mcp",
            "version": "0.9.0",
            "seenAt": "2026-10-04T16:38:03.126836564Z"
          },
          {
            "registry": "pypi",
            "name": "replicate",
            "version": "1.0.7",
            "released": "2025-05-27",
            "seenAt": "2026-10-04T16:38:01.981507625Z"
          }
        ],
        "githubStars": 9486,
        "npmWeekly": 705916,
        "pypiWeekly": 374867,
        "securityTxt": {
          "url": "https://replicate.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:45.229571506Z"
        },
        "llmsTxt": {
          "url": "https://replicate.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:18:09.902788266Z"
        },
        "domain": {
          "domain": "replicate.com",
          "registered": "1998-05-26",
          "source": "https://rdap.verisign.com/com/v1/domain/replicate.com",
          "checkedAt": "2026-10-04T13:07:04.742407865Z"
        },
        "pages": [
          {
            "url": "https://replicate.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-04T15:47:17.68615995Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "490f4836aca3"
          },
          {
            "url": "https://replicate.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:19.974502738Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3f1305f154be"
          },
          {
            "url": "https://replicate.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:22.274270654Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "8e299fbc64eb"
          },
          {
            "url": "https://replicate.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:23.877388602Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ea48efe3382b"
          }
        ],
        "updatedAt": "2026-10-04T22:35:30.021478789Z"
      }
    },
    "summary": "Replicate Deployments has a score of 63.7 (B) against Beam's 55.5 (C). Both do compute gpu. The largest gap is schema \u0026 documentation, 27 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy",
    "json": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.md",
    "slim": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.min.md"
  },
  "markdown": "Replicate Deployments has a score of 63.7 (B) against Beam's 55.5 (C). Both do compute gpu. The largest gap is schema \u0026 documentation, 27 points.\n\n- Beam: grade C, 55.5/100, rank #313 of 452. Markdown https://www.anchorterminal.com/tools/beam.md · JSON https://www.anchorterminal.com/api/v1/tools/beam.json\n- Replicate Deployments: grade B, 63.7/100, rank #197 of 452. Markdown https://www.anchorterminal.com/tools/replicate-deploy.md · JSON https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Which one, for what\n\nPick Beam for security \u0026 auth (+10), payments \u0026 pricing (+10), maintenance \u0026 community (+15).\n\nPick Replicate Deployments for reliability (+20), schema \u0026 documentation (+27), agent ergonomics (+10), transparency \u0026 trust (+6).\n\n## Score by category\n\n| Category | Weight | Beam | Replicate Deployments | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 55 | 75 | Replicate Deployments +20 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 58 | 85 | Replicate Deployments +27 |\n| Agent ergonomics | 13% (16.2 this run) | 58 | 68 | Replicate Deployments +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 50 | 40 | Beam +10 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 30 | Beam +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 85 | 70 | Beam +15 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 74 | 80 | Replicate Deployments +6 |\n| Negative events | ≤15 | -2 | 0 | |\n| **Total** | | **55.5 · C** | **63.7 · B** | |\n\n## Facts side by side\n\n| Fact | Beam | Replicate Deployments |\n| --- | --- | --- |\n| Kind | Model platform | HTTP API |\n| Vendor | Beam | Replicate |\n| Hosted endpoint | `https://app.beam.cloud/api/v1` | `https://api.replicate.com/v1` |\n| Transports | HTTP | HTTP, SSE (legacy), stdio |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | AGPL-3.0 | Apache-2.0 |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| MCP registry | not listed | not listed |\n| Last release | 2026-10-01 | 2026-09-22 |\n| Popularity | 1.8k stars, 8.5k PyPI/wk | 9.5k stars, 634k npm/wk, 387k PyPI/wk |\n| Agent reviews | 3/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**Beam.** Per-millisecond billing with cold starts and image pulls free, H100 PCIe at $3.50 and RTX 4090 at $0.69 an hour. No published request rate limits, 429 handling or SLA.\n\n**Replicate Deployments.** OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.\n\n## Before you call either\n\n### Beam\n\n1. Check the response body for `ok: false` on gateway calls; a failure can arrive as HTTP 200\n2. Don't pipe `beam deploy --format json` output into CI logs, since it contains the workspace token\n3. Route anything over 180 seconds to a task queue and poll the task instead of holding the endpoint request\n4. Set `keep_warm_seconds` deliberately; the 180-second endpoint default bills three minutes of GPU after every call\n5. Pass `gpu=[\"RTX4090\", \"A10G\"]` so a job still schedules when one type is out\n\n### Replicate Deployments\n\n1. List `GET /v1/hardware` first and use the returned `sku` in the deployment body\n2. Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not\n3. Send `Prefer: wait` on deployment predictions to block instead of polling\n4. Copy outputs within an hour; API prediction data is deleted after that\n5. Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute\n\n## Other comparisons with Beam or Replicate Deployments\n\n- [Baseten vs Beam](https://www.anchorterminal.com/compare/baseten-vs-beam.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Beam vs Koyeb](https://www.anchorterminal.com/compare/beam-vs-koyeb.md)\n- [Beam vs Lambda Cloud](https://www.anchorterminal.com/compare/beam-vs-lambda.md)\n- [Beam vs Modal](https://www.anchorterminal.com/compare/beam-vs-modal.md)\n- [Beam vs Northflank](https://www.anchorterminal.com/compare/beam-vs-northflank.md)\n- [Beam vs Runpod](https://www.anchorterminal.com/compare/beam-vs-runpod.md)\n- [Koyeb vs Replicate Deployments](https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.md)\n- [Lambda Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md)\n- [Modal vs Replicate Deployments](https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.md)\n- [Northflank vs Replicate Deployments](https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.md)\n- [Replicate Deployments vs Runpod](https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Beam vs Replicate Deployments",
        "url": ""
      }
    ],
    "description": "Replicate Deployments has a score of 63.7 (B) against Beam's 55.5 (C). Both do compute gpu. The largest gap is schema \u0026 documentation, 27 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Beam C 55.5",
      "Replicate Deployments B 63.7",
      "scores"
    ],
    "h1": "Beam vs Replicate Deployments",
    "image": "https://www.anchorterminal.com/assets/og/compare-beam-vs-replicate-deploy.png",
    "path": "/compare/beam-vs-replicate-deploy",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Beam vs Replicate Deployments for AI agents, C 55.5 vs B 63.7",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy"
  },
  "tokens": {
    "markdown": 1450,
    "slim": 330
  },
  "version": 1
}
