{
  "data": {
    "a": {
      "slug": "cloudflare-workers-ai",
      "name": "Cloudflare Workers AI",
      "vendor": "Cloudflare, Inc.",
      "vendorUrl": "https://www.cloudflare.com",
      "kind": "model",
      "category": "inference",
      "summary": "Workers AI is Cloudflare's serverless inference service for open-weight models, covering text generation, embeddings, images and speech. Agents call it through the Cloudflare REST API, OpenAI-compatible endpoints or a Workers binding.",
      "url": "https://www.anchorterminal.com/tools/cloudflare-workers-ai",
      "markdownUrl": "https://www.anchorterminal.com/tools/cloudflare-workers-ai.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cloudflare-workers-ai.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cloudflare-workers-ai.json",
      "repo": "https://github.com/cloudflare/api-schemas",
      "license": "Proprietary service under Cloudflare's Self-Serve Subscription Agreement. Each hosted model carries its own open-weight licence, linked from its model page. The `cloudflare` SDKs are Apache-2.0, `workers-ai-provider` is MIT and the OpenAPI repository is BSD-3-Clause",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
      "packages": [
        {
          "registry": "npm",
          "name": "workers-ai-provider"
        },
        {
          "registry": "npm",
          "name": "cloudflare"
        },
        {
          "registry": "pypi",
          "name": "cloudflare"
        }
      ],
      "auth": "api-key",
      "authNotes": "A Cloudflare API token as a Bearer header, created by a person in the dashboard, with the Workers AI Read or Workers AI Edit permission on an account. Tokens can carry an expiry and an IP address filter. Inside a Worker the `AI` binding needs no token. The API reference also accepts the older global API key with an account email. Later tokens can be created through the API with an existing token.",
      "pricing": "freemium",
      "pricingNotes": "$0.011 per 1,000 neurons, shown per model as token, image or audio prices, for example Gemma 4 26B at $0.10 in and $0.30 out per 1M tokens. 10,000 neurons a day are free on the Workers Free and Paid plans, and going past that needs Workers Paid, from $5 a month. Seven frontier models need Workers Paid or prepaid AI Gateway credits. A Free plan account starts without a contract (https://developers.cloudflare.com/workers-ai/platform/pricing/, checked 2026-10-09).",
      "priceSummary": "from $0.0605 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "partial",
        "evidence": "AI Gateway Machine Payments, in beta since 30 September 2026, takes x402 payment on `POST /ai/run` for four open models when the request carries `Payment-Method: x402`. It needs a Cloudflare API token, a United States account and a card on file. The per-model route `/ai/run/{model_name}` and the OpenAI-style routes are not covered. We did not call it (https://developers.cloudflare.com/ai-gateway/features/machine-payments/, checked 2026-10-09).",
        "endpoints": [
          {
            "url": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai/run",
            "priceUsd": null,
            "network": ""
          }
        ]
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 379824,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.cloudflare.com/workers-ai/",
      "rateLimitsUrl": "https://developers.cloudflare.com/workers-ai/platform/limits/",
      "llmsTxt": "https://developers.cloudflare.com/workers-ai/llms.txt",
      "openapi": "https://raw.githubusercontent.com/cloudflare/api-schemas/main/openapi.json",
      "capabilities": [
        "inference.llm",
        "inference.open-weights",
        "embed.text",
        "rerank",
        "image.generate",
        "speech.stt",
        "speech.tts",
        "compute.batch"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "usage-based",
        "free-tier",
        "x402",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "batch",
        "prompt-caching",
        "typescript",
        "python",
        "go",
        "status-page",
        "bug-bounty"
      ],
      "lastRelease": "2026-10-01",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 68.3,
        "grade": "B",
        "agentReady": false,
        "rank": 224,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 7,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 72,
          "payments": 52,
          "reliability": 66,
          "schema": 80,
          "security": 77,
          "transparency": 76
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": -3,
        "negativeNotes": [
          "28 July 2026. Kimi K2.6, Kimi K2.7 Code and GLM-5.2 began returning 403 (code 5035) on the Workers Free plan on the day of the notice, with no advance warning found. Three points (https://developers.cloudflare.com/changelog/post/2026-07-28-models-require-workers-paid/)."
        ],
        "verdict": "Per-model prices, rate limits and JSON Schemas are public, 10,000 neurons a day are free, and x402 payment is in beta on `/ai/run` for four models. No SLA was found, three models moved to paid-only access on 28 July 2026 with no notice, and several docs pages still use a model retired in May.",
        "bestFor": "Agents that already run on Cloudflare Workers, or that want open-weight chat, embedding, reranking, speech and image models behind one token with a free daily allowance.",
        "strengths": [
          "Per-model token prices, cached-input prices and neuron rates are published without a login, with 10,000 neurons a day free on the Workers Free plan",
          "x402 payment (Machine Payments, beta since 30 September 2026) works on `POST /ai/run` for four open models",
          "Rate limits are published per task type, and a 429 separates a spent free allocation (3036) from capacity (3040)",
          "API tokens are limited to Workers AI permissions on an account, with optional expiry and IP address filtering",
          "DeepSeek V4 and GLM-5.3 run with a 1,048,576-token context, with prefix caching and discounted cached input"
        ],
        "weaknesses": [
          "No SLA for Workers AI was found in the docs or the agreements read",
          "Three models moved to paid-only access on 28 July 2026, the day of the notice, and Free plan calls to them return 403",
          "Eighteen models were retired on 30 May 2026 with 22 days' notice, and no minimum notice period is published",
          "The REST quick start, the OpenAI-compatibility page and the JSON Mode model list still name models on the 30 May 2026 retirement list",
          "The x402 route still needs a Cloudflare API token, a United States account and a card on file"
        ],
        "agentNotes": [
          "Check the model page before calling. Seven models need Workers Paid or AI Gateway credits and return 403 with code 5035 on the Free plan",
          "Read the internal code on a 429. 3036 means the day's 10,000 free neurons are spent until 00:00 UTC, 3040 means capacity, so retry later",
          "Set `options.rejectIfBusy` to fail fast, or add `?queueRequest=true` for the batch route when the answer can wait",
          "Send the same `x-session-affinity` value on every turn of a session to reach the prefix cache and the cached-input price",
          "Keep paid frontier models under 20 requests a minute per model, or 50 with prepaid AI Gateway credits"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 68.3
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 72,
          "payments": 52,
          "reliability": 66,
          "schema": 80,
          "security": 77,
          "transparency": 56
        },
        "provenanceScore": 95
      },
      "connect": {
        "http": "curl https://api.cloudflare.com/client/v4/accounts/$CLOUDFLARE_ACCOUNT_ID/ai/run/@cf/google/gemma-4-26b-a4b-it \\\n  -X POST \\\n  -H \"Authorization: Bearer $CLOUDFLARE_AUTH_TOKEN\" \\\n  -d '{ \"messages\": [{ \"role\": \"system\", \"content\": \"You are a friendly assistant\" }, { \"role\": \"user\", \"content\": \"Why is pizza so good\" }]}'",
        "x402": "curl -iX POST \"https://api.cloudflare.com/client/v4/accounts/$CLOUDFLARE_ACCOUNT_ID/ai/run\" \\\n  --header \"Authorization: Bearer $CLOUDFLARE_API_TOKEN\" \\\n  --header \"Payment-Method: x402\" \\\n  --header \"Content-Type: application/json\" \\\n  --data '{\"model\":\"z-ai/glm-4.7-flash\",\"input\":{\"messages\":[{\"role\":\"user\",\"content\":\"What is Cloudflare?\"}]}}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/cloudflare-workers-ai"
      },
      "sameCompany": [
        "cloudflare-ai-gateway",
        "cloudflare-sandbox-sdk",
        "cloudflare-web-search",
        "cloudflare-mcp",
        "cloudflare-email-service",
        "cloudflare-r2",
        "cloudflare-clef"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "Whisper large v3 turbo, speech to text",
          "unit": "audio-minute",
          "usd": 0.0005,
          "note": "46.63 neurons"
        },
        {
          "item": "Deepgram Nova-3, speech to text",
          "unit": "audio-minute",
          "usd": 0.0052,
          "note": "472.73 neurons. $0.0092 over WebSocket"
        },
        {
          "item": "Deepgram Aura-1, text to speech",
          "unit": "1m-chars",
          "usd": 15,
          "note": "listed as $0.015 per 1,000 characters"
        },
        {
          "item": "Deepgram Aura-2, text to speech",
          "unit": "1m-chars",
          "usd": 30,
          "note": "listed as $0.030 per 1,000 characters, English and Spanish"
        }
      ],
      "provenance": {
        "legalEntity": "Cloudflare, Inc.",
        "domain": "cloudflare.com",
        "domainRegistered": "2009-02-17",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cloudflare.com/terms/",
        "privacy": "https://www.cloudflare.com/privacypolicy/",
        "statusPage": "https://www.cloudflarestatus.com",
        "changelog": "https://developers.cloudflare.com/changelog/product/workers-ai/",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "notes": [
          "Terms are the Self-Serve Subscription Agreement, last updated 12 September 2025, which names Cloudflare, Inc. at 101 Townsend St, San Francisco. The Workers AI data usage page cites it and the Enterprise Subscription Agreement as the governing documents.",
          "The service-specific terms (https://www.cloudflare.com/service-specific-terms-developer-platform/), last updated 28 September 2026, have a section for Workers AI and AI Gateway that supplements the agreement.",
          "The privacy policy is effective from 4 November 2025.",
          "www.cloudflare.com/.well-known/security.txt, read on 9 October 2026, names HackerOne and a disclosure policy and has no Expires field. It is recorded as valid to match the other Cloudflare listings.",
          "The endpoint is on api.cloudflare.com.",
          "The status page lists Workers AI as its own component, operational on 9 October 2026.",
          "RDAP for cloudflare.com gives a registration date of 2009-02-17.",
          "The Workers AI changelog page (https://developers.cloudflare.com/workers-ai/changelog/) stops at 16 June 2026, so the product changelog is recorded."
        ],
        "score": 95
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cloudflare-workers-ai.json",
      "live": {
        "slug": "cloudflare-workers-ai",
        "probe": {
          "target": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
          "method": "get",
          "lastAt": "2026-10-10T02:54:51.330012119Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 19,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 20,
          "p95ms24h": 49,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://www.cloudflarestatus.com",
          "indicator": "major",
          "summary": "Partial System Outage",
          "checkedAt": "2026-10-10T02:50:03.008405875Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "cloudflare",
            "version": "7.3.0",
            "seenAt": "2026-10-09T16:46:16.222261099Z"
          },
          {
            "registry": "npm",
            "name": "workers-ai-provider",
            "version": "4.0.0",
            "seenAt": "2026-10-09T16:46:15.310826363Z"
          },
          {
            "registry": "pypi",
            "name": "cloudflare",
            "version": "5.9.0",
            "released": "2026-10-03",
            "seenAt": "2026-10-09T16:46:17.361988865Z"
          }
        ],
        "githubStars": 194,
        "npmWeekly": 379824,
        "pypiWeekly": 365643,
        "pages": [
          {
            "url": "https://developers.cloudflare.com/changelog/product/workers-ai/",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:50.279057567Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "486cf84220c1"
          },
          {
            "url": "https://developers.cloudflare.com/changelog/post/2026-05-08-planned-model-deprecations/",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:42.235428702Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "456d5c31e53a"
          },
          {
            "url": "https://developers.cloudflare.com/changelog/post/2026-07-28-models-require-workers-paid/",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:44.226321807Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "23663c817503"
          }
        ],
        "updatedAt": "2026-10-10T02:54:51.330012119Z"
      }
    },
    "answer": "Cloudflare Workers AI scores 68.3 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories.",
    "b": {
      "slug": "prism-inference",
      "name": "Prism Inference",
      "vendor": "Prism Technologies Inc",
      "vendorUrl": "https://prisminference.com",
      "kind": "model",
      "category": "inference",
      "summary": "Prism is a hosted inference API from Prism Technologies Inc for open-weight models, aimed at coding agents. It accepts OpenAI Chat Completions, OpenAI Responses and Anthropic Messages requests at api.prisminference.com. It launched on 24 September 2026.",
      "url": "https://www.anchorterminal.com/tools/prism-inference",
      "markdownUrl": "https://www.anchorterminal.com/tools/prism-inference.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/prism-inference.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/prism-inference.json",
      "repo": "https://github.com/prismhq/hermes-prism-provider",
      "license": "Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.prisminference.com/v1",
      "packages": [],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` with a key issued from account settings after sign-up at prisminference.com/signup. The inference endpoints also accept the key in `x-api-key`. A missing, invalid, expired or revoked key returns 401. No key scopes were found in the reviewed documentation. An agent can call `POST https://prisminference.com/api/agent-signups` with the owner's email and a username and receive a key once, but that key can't run inference until the owner supplies a six-digit emailed code, and it expires after 30 days. `GET /v1/models` needs no key.",
      "pricing": "usage",
      "pricingNotes": "Prepaid per-token pricing with no minimum. DeepSeek-V4.1-Flash is $0.09 in, $1.20 out and $0.06 cache read per million tokens, and Gemma 4 31B is $0.30, $0.40 and $0.15 (https://prisminference.com/pricing, matched by https://api.prisminference.com/v1/models). No free tier or trial credit was found, and a workspace without credit gets 402. A person funds the workspace in a browser. Elastic endpoints, dedicated deployments and batch are sold through sales with no published price.",
      "priceSummary": "from $0.09 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in llms.txt, the docs index, the OpenAPI file or the pricing page, read 2026-10-08. The docs describe prepaid credit funded by a person in the browser.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.prisminference.com",
      "rateLimitsUrl": "https://docs.prisminference.com/rate-limits",
      "llmsTxt": "https://prisminference.com/llms.txt",
      "openapi": "https://docs.prisminference.com/openapi.yaml",
      "capabilities": [
        "inference.fast",
        "inference.open-weights",
        "inference.llm"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "fast",
        "usage-priced",
        "prepaid",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "anthropic-compatible",
        "zero-retention",
        "status-page",
        "new"
      ],
      "lastRelease": "2026-10-06",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 60.1,
        "grade": "C",
        "agentReady": false,
        "rank": 528,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 13,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 49,
          "payments": 30,
          "reliability": 65,
          "schema": 82,
          "security": 65,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -2,
        "negativeNotes": [
          "2026-10-08. The home page shows a '99.99% Uptime SLA' tile, and the pricing page says 'No minimums, no rate limits' above the per-token table. The terms of 9 September 2026 say the services have no guaranteed uptime or service credit unless a separate written agreement says otherwise, and the docs describe per-key rate limits that return 429. No SLA document was found. A misleading claim, with the smallest deduction because the terms and docs state the real position (https://prisminference.com/, https://prisminference.com/pricing, https://prisminference.com/terms, https://docs.prisminference.com/rate-limits)."
        ],
        "verdict": "Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found.",
        "bestFor": "Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.",
        "strengths": [
          "One key works across OpenAI Chat Completions, OpenAI Responses and Anthropic Messages, with a public OpenAPI 3.1 file, llms.txt and Markdown docs",
          "Zero data retention is the default on every tier, and the privacy policy, terms and docs all say inputs and outputs are never used for training",
          "Every error carries a stable `code`, a `retryable` flag, a `fix` hint and a `docs_url`, and 429 carries `Retry-After` in seconds",
          "`GET /v1/models` answers without a key and returns context length, maximum output and per-token prices for each model",
          "DeepSeek-V4.1-Flash is listed with a 1M-token context and 384,000 output tokens at $0.09 in and $1.20 out per million"
        ],
        "weaknesses": [
          "Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available",
          "No rate-limit numbers are published. The docs say per-key limits exist, and the pricing page says 'no rate limits'",
          "The home page shows a '99.99% Uptime SLA' tile, while the terms say there is no guaranteed uptime without a separate written agreement",
          "No free tier found. Billing is prepaid, a person funds the workspace in a browser, and the agent sign-up key needs an emailed code",
          "The service launched on 24 September 2026. The changelog has one entry, and no deprecation policy, sub-processor list or certification was found"
        ],
        "agentNotes": [
          "Call `GET https://api.prisminference.com/v1/models` at start-up, with no key, and use only ids it returns. Expect 403 on `gemma-4-31b` without organisation access",
          "Use base URL `https://api.prisminference.com/v1` for OpenAI clients and `https://api.prisminference.com` with no `/v1` for Anthropic clients",
          "Read `error.retryable` before retrying, and wait for `Retry-After` on 429, which covers both key limits and model capacity",
          "Send `reasoning_effort: \"none\"` or `low` when latency matters. Reasoning is on by default and its tokens are billed as output",
          "Keep conversation state yourself and send `store: false` on Responses. `previous_response_id`, stored responses and hosted tools aren't supported"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 60.1
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 49,
          "payments": 30,
          "reliability": 65,
          "schema": 82,
          "security": 65,
          "transparency": 43
        },
        "provenanceScore": 79
      },
      "connect": {
        "install": "pip install openai   # or: npm install openai, base URL https://api.prisminference.com/v1. Anthropic SDKs use https://api.prisminference.com with no /v1",
        "http": "curl \"https://api.prisminference.com/v1/chat/completions\" \\\n  -H \"Authorization: Bearer $PRISM_API_KEY\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\"model\":\"deepseek-v4.1-flash\",\"messages\":[{\"role\":\"user\",\"content\":\"Return pong.\"}]}'",
        "claudeCode": "export ANTHROPIC_BASE_URL=https://api.prisminference.com\nexport ANTHROPIC_AUTH_TOKEN=$PRISM_API_KEY\nexport ANTHROPIC_MODEL=deepseek-v4.1-flash\nexport ANTHROPIC_SMALL_FAST_MODEL=gemma-4-31b"
      },
      "letme": {
        "capability": "https://letme.dev/inference.fast",
        "tool": "https://letme.dev/prism-inference"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "Prism Technologies Inc",
        "domain": "prisminference.com",
        "domainRegistered": "2026-09-09",
        "endpointOnVendorDomain": true,
        "terms": "https://prisminference.com/terms",
        "privacy": "https://prisminference.com/privacy",
        "statusPage": "https://status.prisminference.com",
        "changelog": "https://docs.prisminference.com/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "notes": [
          "The terms and the privacy policy, both last updated 9 September 2026, name Prism Technologies Inc. The terms are governed by California law with arbitration in San Francisco.",
          "RDAP gives a registration date of 2026-09-09 for prisminference.com, with Name.com as registrar.",
          "security.txt has a Contact line (founders@prisminference.com), an Expires date of 2027-10-06 and a Canonical line. No disclosure policy or bug bounty is named.",
          "The status page runs on incident.io with one component for each model and no component for the API or the website. Its incidents feed was empty on 8 October 2026.",
          "The YC directory lists Prism in the Spring 2025 batch, in San Francisco, with a team of 2. The home page links an X account named prism_videos, from the company's earlier video product."
        ],
        "score": 79
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/prism-inference.json",
      "live": {
        "slug": "prism-inference",
        "probe": {
          "target": "https://api.prisminference.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:06.706259541Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 131,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 149,
          "p95ms24h": 452,
          "samples24h": 249,
          "samples30d": 373,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 93,
              "ok": 93
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.prisminference.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T02:50:43.671876128Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "prismhq/hermes-prism-provider",
            "version": "v1.0.2",
            "released": "2026-09-16",
            "seenAt": "2026-10-09T17:14:33.431257688Z"
          }
        ],
        "githubStars": 0,
        "securityTxt": {
          "url": "https://prisminference.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-10-06T00:00:00.000Z",
          "checkedAt": "2026-10-09T15:40:35.003140299Z"
        },
        "llmsTxt": {
          "url": "https://prisminference.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:37.937240222Z"
        },
        "pages": [
          {
            "url": "https://docs.prisminference.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-09T18:38:03.252133069Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "184a0fafdb8c"
          },
          {
            "url": "https://prisminference.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:52.542857787Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "dffb5d16a3c8"
          },
          {
            "url": "https://prisminference.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:54.77090078Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "64ea7b8f1b7b"
          },
          {
            "url": "https://prisminference.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:56.713185583Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "27dd15c69f49"
          }
        ],
        "updatedAt": "2026-10-10T02:55:06.706259541Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Cloudflare, Inc.",
        "b": "Prism Technologies Inc",
        "name": "Vendor"
      },
      {
        "a": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
        "b": "https://api.prisminference.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "payer tooling only",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cloudflare's Self-Serve Subscription Agreement. Each hosted model carries its own open-weight licence, linked from its model page. The `cloudflare` SDKs are Apache-2.0, `workers-ai-provider` is MIT and the OpenAPI repository is BSD-3-Clause",
        "b": "Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-10-01",
        "b": "2026-10-06",
        "name": "Last release"
      },
      {
        "a": "2025-09-12",
        "b": "2026-09-09",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2026-09-09",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "380k npm/wk",
        "b": "none",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "Cloudflare Workers AI scores 68.3 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories.",
        "question": "Which is better for AI agents, Cloudflare Workers AI or Prism Inference?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Cloudflare Workers AI and Prism Inference need an API key?"
      },
      {
        "answer": "Yes. Cloudflare Workers AI has a hosted endpoint at https://api.cloudflare.com/client/v4/accounts/{account_id}/ai and Prism Inference at https://api.prisminference.com/v1.",
        "question": "Can an agent call Cloudflare Workers AI and Prism Inference without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Agent ergonomics, 75 against 68",
          "Security \u0026 auth, 77 against 65",
          "Payments \u0026 pricing, 52 against 30",
          "Maintenance \u0026 community, 72 against 49",
          "Transparency \u0026 trust, 76 against 61"
        ],
        "also": null,
        "goodFor": "Agents that already run on Cloudflare Workers, or that want open-weight chat, embedding, reranking, speech and image models behind one token with a free daily allowance.",
        "slug": "cloudflare-workers-ai",
        "watchFor": "No SLA for Workers AI was found in the docs or the agreements read"
      },
      {
        "aheadOn": null,
        "also": null,
        "goodFor": "Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.",
        "slug": "prism-inference",
        "watchFor": "Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available"
      }
    ],
    "job": {
      "capability": "inference.llm",
      "name": "LLM inference"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-workers-ai.json",
        "title": "Claude API vs Cloudflare Workers AI",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-workers-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference.json",
        "title": "Claude API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-cloudflare-workers-ai.json",
        "title": "Antseed vs Cloudflare Workers AI",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-cloudflare-workers-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-prism-inference.json",
        "title": "Antseed vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-workers-ai.json",
        "title": "BlockRun.AI vs Cloudflare Workers AI",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-workers-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference.json",
        "title": "BlockRun.AI vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai.json",
        "title": "Cloudflare AI Gateway vs Cloudflare Workers AI",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.json",
        "title": "Cloudflare AI Gateway vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-cohere-chat.json",
        "title": "Cloudflare Workers AI vs Cohere Chat API (Command models)",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-cohere-chat"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepinfra.json",
        "title": "Cloudflare Workers AI vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepseek-api.json",
        "title": "Cloudflare Workers AI vs DeepSeek API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepseek-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-gemini-api.json",
        "title": "Cloudflare Workers AI vs Gemini Developer API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-gemini-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-groq.json",
        "title": "Cloudflare Workers AI vs GroqCloud",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-groq"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-mistral-api.json",
        "title": "Cloudflare Workers AI vs Mistral AI API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-mistral-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-novita-ai.json",
        "title": "Cloudflare Workers AI vs Novita AI",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-novita-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openai-api.json",
        "title": "Cloudflare Workers AI vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openrouter.json",
        "title": "Cloudflare Workers AI vs OpenRouter",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-sambanova.json",
        "title": "Cloudflare Workers AI vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-siliconflow.json",
        "title": "Cloudflare Workers AI vs SiliconFlow",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-siliconflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference.json",
        "title": "Cohere Chat API (Command models) vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.json",
        "title": "DeepInfra vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference.json",
        "title": "DeepSeek API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference.json",
        "title": "Gemini Developer API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-vs-prism-inference.json",
        "title": "GroqCloud vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/groq-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference.json",
        "title": "Mistral AI API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference.json",
        "title": "Novita AI vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.json",
        "title": "OpenAI API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openrouter-vs-prism-inference.json",
        "title": "OpenRouter vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/openrouter-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow.json",
        "title": "Prism Inference vs SiliconFlow",
        "url": "https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.json",
        "title": "Prism Inference vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova"
      }
    ],
    "scores": [
      {
        "by": 1,
        "cloudflare-workers-ai": 66,
        "edge": "cloudflare-workers-ai",
        "key": "reliability",
        "name": "Reliability",
        "prism-inference": 65,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 2,
        "cloudflare-workers-ai": 80,
        "edge": "prism-inference",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "prism-inference": 82,
        "weight": 13
      },
      {
        "by": 7,
        "cloudflare-workers-ai": 75,
        "edge": "cloudflare-workers-ai",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "prism-inference": 68,
        "weight": 13
      },
      {
        "by": 12,
        "cloudflare-workers-ai": 77,
        "edge": "cloudflare-workers-ai",
        "key": "security",
        "name": "Security \u0026 auth",
        "prism-inference": 65,
        "weight": 14
      },
      {
        "by": 22,
        "cloudflare-workers-ai": 52,
        "edge": "cloudflare-workers-ai",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "prism-inference": 30,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 23,
        "cloudflare-workers-ai": 72,
        "edge": "cloudflare-workers-ai",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "prism-inference": 49,
        "weight": 7
      },
      {
        "by": 15,
        "cloudflare-workers-ai": 76,
        "edge": "cloudflare-workers-ai",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "prism-inference": 61,
        "weight": 7
      }
    ],
    "summary": "Cloudflare Workers AI scores 68.3 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Both do llm inference.",
    "verdicts": {
      "cloudflare-workers-ai": "Per-model prices, rate limits and JSON Schemas are public, 10,000 neurons a day are free, and x402 payment is in beta on `/ai/run` for four models. No SLA was found, three models moved to paid-only access on 28 July 2026 with no notice, and several docs pages still use a model retired in May.",
      "prism-inference": "Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference",
    "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.md",
    "slim": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.min.md"
  },
  "markdown": "Cloudflare Workers AI scores 68.3 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Both do llm inference.\n\n- Cloudflare Workers AI: grade B, 68.3/100, rank #224 of 950. Markdown https://www.anchorterminal.com/tools/cloudflare-workers-ai.md · JSON https://www.anchorterminal.com/api/v1/tools/cloudflare-workers-ai.json\n- Prism Inference: grade C, 60.1/100, rank #528 of 950. Markdown https://www.anchorterminal.com/tools/prism-inference.md · JSON https://www.anchorterminal.com/api/v1/tools/prism-inference.json\n- Best model APIs and inference for AI agents: https://www.anchorterminal.com/best/inference/index.md\n- All 136 models comparisons: https://www.anchorterminal.com/compare/inference/index.md\n\n## Which one, for what\n\n### Cloudflare Workers AI (B)\n\nGood for: Agents that already run on Cloudflare Workers, or that want open-weight chat, embedding, reranking, speech and image models behind one token with a free daily allowance.\n\nAhead on:\n- Agent ergonomics, 75 against 68\n- Security \u0026 auth, 77 against 65\n- Payments \u0026 pricing, 52 against 30\n- Maintenance \u0026 community, 72 against 49\n- Transparency \u0026 trust, 76 against 61\n\nWatch for: No SLA for Workers AI was found in the docs or the agreements read\n\n### Prism Inference (C)\n\nGood for: Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.\n\nWatch for: Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available\n\n\n## Score by category\n\n| Category | Weight | Cloudflare Workers AI | Prism Inference | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 66 | 65 | Cloudflare Workers AI +1 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 80 | 82 | Prism Inference +2 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 68 | Cloudflare Workers AI +7 |\n| Security \u0026 auth | 14% (17.5 this run) | 77 | 65 | Cloudflare Workers AI +12 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 52 | 30 | Cloudflare Workers AI +22 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 72 | 49 | Cloudflare Workers AI +23 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 76 | 61 | Cloudflare Workers AI +15 |\n| Negative events | ≤15 | -3 | -2 | |\n| **Total** | | **68.3 · B** | **60.1 · C** | |\n\n## Facts side by side\n\n| Fact | Cloudflare Workers AI | Prism Inference |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Cloudflare, Inc. | Prism Technologies Inc |\n| Hosted endpoint | `https://api.cloudflare.com/client/v4/accounts/{account_id}/ai` | `https://api.prisminference.com/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | payer tooling only | no |\n| Licence | Proprietary service under Cloudflare's Self-Serve Subscription Agreement. Each hosted model carries its own open-weight licence, linked from its model page. The `cloudflare` SDKs are Apache-2.0, `workers-ai-provider` is MIT and the OpenAPI repository is BSD-3-Clause | Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-10-01 | 2026-10-06 |\n| Terms last updated | 2025-09-12 | 2026-09-09 |\n| Privacy policy last updated | no date given | 2026-09-09 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | yes |\n| Terms restrict benchmarking | not found in the text | not found in the text |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | yes | yes |\n| Popularity | 380k npm/wk | none |\n\n## Verdicts\n\n**Cloudflare Workers AI.** Per-model prices, rate limits and JSON Schemas are public, 10,000 neurons a day are free, and x402 payment is in beta on `/ai/run` for four models. No SLA was found, three models moved to paid-only access on 28 July 2026 with no notice, and several docs pages still use a model retired in May.\n\n**Prism Inference.** Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found.\n\n## Before you call either\n\n### Cloudflare Workers AI\n\n1. Check the model page before calling. Seven models need Workers Paid or AI Gateway credits and return 403 with code 5035 on the Free plan\n2. Read the internal code on a 429. 3036 means the day's 10,000 free neurons are spent until 00:00 UTC, 3040 means capacity, so retry later\n3. Set `options.rejectIfBusy` to fail fast, or add `?queueRequest=true` for the batch route when the answer can wait\n4. Send the same `x-session-affinity` value on every turn of a session to reach the prefix cache and the cached-input price\n5. Keep paid frontier models under 20 requests a minute per model, or 50 with prepaid AI Gateway credits\n\n### Prism Inference\n\n1. Call `GET https://api.prisminference.com/v1/models` at start-up, with no key, and use only ids it returns. Expect 403 on `gemma-4-31b` without organisation access\n2. Use base URL `https://api.prisminference.com/v1` for OpenAI clients and `https://api.prisminference.com` with no `/v1` for Anthropic clients\n3. Read `error.retryable` before retrying, and wait for `Retry-After` on 429, which covers both key limits and model capacity\n4. Send `reasoning_effort: \"none\"` or `low` when latency matters. Reasoning is on by default and its tokens are billed as output\n5. Keep conversation state yourself and send `store: false` on Responses. `previous_response_id`, stored responses and hosted tools aren't supported\n\n## Questions\n\n### Which is better for AI agents, Cloudflare Workers AI or Prism Inference?\n\nCloudflare Workers AI scores 68.3 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories.\n\n### Do Cloudflare Workers AI and Prism Inference need an API key?\n\nBoth need an API key.\n\n### Can an agent call Cloudflare Workers AI and Prism Inference without installing anything?\n\nYes. Cloudflare Workers AI has a hosted endpoint at https://api.cloudflare.com/client/v4/accounts/{account_id}/ai and Prism Inference at https://api.prisminference.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cloudflare-workers-ai\", \"b\": \"prism-inference\"}`. From a terminal: `anchor compare cloudflare-workers-ai prism-inference`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cloudflare-workers-ai.json and https://www.anchorterminal.com/api/v1/tools/prism-inference.json\n\n## Other comparisons with Cloudflare Workers AI or Prism Inference\n\n- [Claude API vs Cloudflare Workers AI](https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-workers-ai.md)\n- [Claude API vs Prism Inference](https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference.md)\n- [Antseed vs Cloudflare Workers AI](https://www.anchorterminal.com/compare/antseed-vs-cloudflare-workers-ai.md)\n- [Antseed vs Prism Inference](https://www.anchorterminal.com/compare/antseed-vs-prism-inference.md)\n- [BlockRun.AI vs Cloudflare Workers AI](https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-workers-ai.md)\n- [BlockRun.AI vs Prism Inference](https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference.md)\n- [Cloudflare AI Gateway vs Cloudflare Workers AI](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai.md)\n- [Cloudflare AI Gateway vs Prism Inference](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.md)\n- [Cloudflare Workers AI vs Cohere Chat API (Command models)](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-cohere-chat.md)\n- [Cloudflare Workers AI vs DeepInfra](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepinfra.md)\n- [Cloudflare Workers AI vs DeepSeek API](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-deepseek-api.md)\n- [Cloudflare Workers AI vs Gemini Developer API](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-gemini-api.md)\n- [Cloudflare Workers AI vs GroqCloud](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-groq.md)\n- [Cloudflare Workers AI vs Mistral AI API](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-mistral-api.md)\n- [Cloudflare Workers AI vs Novita AI](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-novita-ai.md)\n- [Cloudflare Workers AI vs OpenAI API](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openai-api.md)\n- [Cloudflare Workers AI vs OpenRouter](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-openrouter.md)\n- [Cloudflare Workers AI vs SambaCloud](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-sambanova.md)\n- [Cloudflare Workers AI vs SiliconFlow](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-siliconflow.md)\n- [Cohere Chat API (Command models) vs Prism Inference](https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference.md)\n- [DeepInfra vs Prism Inference](https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.md)\n- [DeepSeek API vs Prism Inference](https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference.md)\n- [Gemini Developer API vs Prism Inference](https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference.md)\n- [GroqCloud vs Prism Inference](https://www.anchorterminal.com/compare/groq-vs-prism-inference.md)\n- [Mistral AI API vs Prism Inference](https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference.md)\n- [Novita AI vs Prism Inference](https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference.md)\n- [OpenAI API vs Prism Inference](https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.md)\n- [OpenRouter vs Prism Inference](https://www.anchorterminal.com/compare/openrouter-vs-prism-inference.md)\n- [Prism Inference vs SiliconFlow](https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow.md)\n- [Prism Inference vs SambaCloud](https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cloudflare Workers AI vs Prism Inference",
        "url": ""
      }
    ],
    "description": "Cloudflare Workers AI scores 68.3 (B) to Prism Inference's 60.1 (C) for llm inference. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Cloudflare Workers AI B 68.3",
      "Prism Inference C 60.1",
      "scores"
    ],
    "h1": "Cloudflare Workers AI vs Prism Inference",
    "image": "https://www.anchorterminal.com/assets/og/compare-cloudflare-workers-ai-vs-prism-inference.png",
    "path": "/compare/cloudflare-workers-ai-vs-prism-inference",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cloudflare Workers AI vs Prism Inference for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference"
  },
  "tokens": {
    "markdown": 2900,
    "slim": 730
  },
  "version": 1
}
