{
  "data": {
    "a": {
      "slug": "cloudflare-ai-gateway",
      "name": "Cloudflare AI Gateway",
      "vendor": "Cloudflare, Inc.",
      "vendorUrl": "https://www.cloudflare.com",
      "kind": "router",
      "category": "inference",
      "summary": "AI Gateway is Cloudflare's proxy and router for model providers. Agents call OpenAI-style, Anthropic-style or native endpoints on the Cloudflare API, with logging, caching, rate and spend limits, retries, fallbacks and one prepaid credit balance.",
      "url": "https://www.anchorterminal.com/tools/cloudflare-ai-gateway",
      "markdownUrl": "https://www.anchorterminal.com/tools/cloudflare-ai-gateway.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cloudflare-ai-gateway.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cloudflare-ai-gateway.json",
      "repo": "https://github.com/cloudflare/ai",
      "license": "Proprietary service under Cloudflare's Self-Serve Subscription Agreement and its service-specific terms. Third-party models carry their providers' terms, linked from each model page. `ai-gateway-provider` is MIT and Cloudflare's OpenAPI repository is BSD-3-Clause",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
      "packages": [
        {
          "registry": "npm",
          "name": "ai-gateway-provider"
        },
        {
          "registry": "npm",
          "name": "cloudflare"
        },
        {
          "registry": "pypi",
          "name": "cloudflare"
        }
      ],
      "auth": "api-key",
      "authNotes": "A Cloudflare API token as a Bearer header, created by a person in the dashboard. The `/ai/*` inference routes need Workers AI Read, and the `/ai-gateway/*` management routes need AI Gateway Read or Edit. Provider-native endpoints on `gateway.ai.cloudflare.com` take the token in `cf-aig-authorization`. Tokens can carry an expiry and an IP address filter and cannot be limited to one gateway. Access is self-serve, with no review or sales step.",
      "pricing": "freemium",
      "pricingNotes": "Free on all plans for analytics, caching and rate limiting, with a Cloudflare account. Third-party models are paid from prepaid Unified Billing credits at provider prices plus a 5 per cent fee on each credit purchase, or with the caller's own provider keys at no Cloudflare charge. Logs for new customers follow Workers Logs pricing, 200,000 events a day free. An account on the Free plan starts without a contract (https://developers.cloudflare.com/ai-gateway/reference/pricing/, checked 2026-10-09).",
      "priceSummary": "5% fee",
      "where": "hosted",
      "x402": {
        "level": "partial",
        "evidence": "Machine Payments, in beta since 30 September 2026, takes x402 payment on `POST /ai/run` for four open models when the request carries `Payment-Method: x402`. It needs a Cloudflare API token, a United States account and a card on file. The OpenAI-style and Anthropic-style routes are not covered. We did not call it (https://developers.cloudflare.com/ai-gateway/features/machine-payments/, checked 2026-10-09).",
        "endpoints": [
          {
            "url": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai/run",
            "priceUsd": null,
            "network": ""
          }
        ]
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 250648,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.cloudflare.com/ai-gateway/",
      "llmsTxt": "https://developers.cloudflare.com/ai-gateway/llms.txt",
      "openapi": "https://raw.githubusercontent.com/cloudflare/api-schemas/main/openapi.json",
      "capabilities": [
        "inference.llm",
        "inference.router",
        "obs.gateway",
        "guard.moderation",
        "guard.pii"
      ],
      "tags": [
        "hosted",
        "router",
        "free-tier",
        "usage-based",
        "x402",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "typescript",
        "python",
        "go",
        "status-page",
        "bug-bounty"
      ],
      "lastRelease": "2026-10-06",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 66.8,
        "grade": "B",
        "agentReady": false,
        "rank": 275,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 73,
          "maintenance": 75,
          "payments": 52,
          "reliability": 66,
          "schema": 75,
          "security": 75,
          "transparency": 73
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": -3,
        "negativeNotes": [
          "6 October 2026. Rejected-credential responses on `POST /ai/run` changed status (ElevenLabs 403, Vertex 500 and other providers' 402 became 401 with code 2009, and Unified Billing rejections became 503) on the day of the changelog entry, with no advance notice found. Three points (https://developers.cloudflare.com/ai-gateway/changelog/)."
        ],
        "verdict": "The gateway layer is free, per-model prices are public with no markup, and logs, spend limits and fallbacks are set per gateway. API tokens cannot be limited to one gateway, the service terms take third-party models outside Cloudflare's DPA, no SLA names the product, and error codes changed on 6 October 2026 with same-day notice.",
        "bestFor": "Teams that already hold a Cloudflare account and want one token, one bill, logs, budgets and fallbacks in front of several model providers, or that want to keep their own provider keys behind a proxy with spend limits.",
        "strengths": [
          "Core gateway functions (analytics, caching, rate limiting) are free on all plans, and provider prices pass through with no markup",
          "One Cloudflare API token reaches 241 catalogue models through OpenAI Chat Completions, OpenAI Responses, Anthropic Messages or the native `/ai/run` route",
          "Spend limits block requests with a 429 once a budget is reached, scoped by model, provider or request metadata, with up to 20 rules a gateway",
          "Logging is switched per gateway and per request, and `cf-aig-collect-log-payload: false` keeps token and cost metadata without storing prompts",
          "x402 payment (Machine Payments, beta since 30 September 2026) works on `POST /ai/run` for four open models"
        ],
        "weaknesses": [
          "API tokens are account-scoped. Any token with AI Gateway Run can send requests through every gateway in the account and spend its stored provider keys",
          "The service-specific terms say the DPA and Information Security Exhibit do not apply to third-party models reached through AI Gateway",
          "No SLA naming AI Gateway was found. The Enterprise SLA lists service-specific SLAs for R2, Queues and Workers only",
          "On 6 October 2026 rejected-credential responses on `/ai/run` changed to 401 with code 2009, or 503 under Unified Billing, announced the same day",
          "Unified Billing traffic is capped at 200 requests per 60 seconds per gateway, and a 5 per cent fee applies to credit purchases",
          "Section 2.2(e) of the Self-Serve Subscription Agreement bars automated agents that generate automated requests into the Services, which matters before any probe is run"
        ],
        "agentNotes": [
          "Use the REST API at `api.cloudflare.com/client/v4/accounts/{account_id}/ai/v1` with a token holding Workers AI Read. A token with only an AI Gateway permission returns 401 with code 10000",
          "Send `cf-aig-gateway-id` on every Workers AI (`@cf/`) call and on dynamic routes. Without it a dynamic route resolves against the default gateway and returns 404",
          "Set `cf-aig-collect-log: false` or `cf-aig-collect-log-payload: false` on sensitive requests. Logging of prompts and responses is on by default",
          "Treat a 429 as a rate limit or a spent budget and a 401 with code 2009 as a rejected provider key. Do not retry a 2009",
          "Set `byok_only` on the gateway or send `cf-aig-no-wholesale: true` to stop a request without a provider key falling through to paid Unified Billing credits"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 66.8
          }
        ],
        "editorialScores": {
          "ergonomics": 73,
          "maintenance": 75,
          "payments": 52,
          "reliability": 66,
          "schema": 75,
          "security": 75,
          "transparency": 51
        },
        "provenanceScore": 95
      },
      "connect": {
        "http": "curl -X POST \"https://api.cloudflare.com/client/v4/accounts/$CLOUDFLARE_ACCOUNT_ID/ai/v1/chat/completions\" \\\n  --header \"Authorization: Bearer $CLOUDFLARE_API_TOKEN\" \\\n  --header \"Content-Type: application/json\" \\\n  --data '{\n    \"model\": \"openai/gpt-4.1-mini\",\n    \"messages\": [{\"role\": \"user\", \"content\": \"What is Cloudflare?\"}]\n  }'",
        "x402": "curl -iX POST \"https://api.cloudflare.com/client/v4/accounts/$CLOUDFLARE_ACCOUNT_ID/ai/run\" \\\n  --header \"Authorization: Bearer $CLOUDFLARE_API_TOKEN\" \\\n  --header \"Payment-Method: x402\" \\\n  --header \"Content-Type: application/json\" \\\n  --data '{\"model\":\"z-ai/glm-4.7-flash\",\"input\":{\"messages\":[{\"role\":\"user\",\"content\":\"What is Cloudflare?\"}]}}'"
      },
      "letme": {
        "capability": "https://letme.dev/inference.llm",
        "tool": "https://letme.dev/cloudflare-ai-gateway"
      },
      "sameCompany": [
        "cloudflare-workers-ai",
        "cloudflare-sandbox-sdk",
        "cloudflare-web-search",
        "cloudflare-mcp",
        "cloudflare-email-service",
        "cloudflare-r2",
        "cloudflare-clef"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "Unified Billing credit purchase fee",
          "unit": "pct",
          "usd": 5,
          "note": "a $100 purchase is charged $105"
        },
        {
          "item": "Logpush requests past 10 million a month",
          "unit": "1k-requests",
          "usd": 0.00005,
          "note": "listed as $0.05 per million, Workers Paid only"
        }
      ],
      "provenance": {
        "legalEntity": "Cloudflare, Inc.",
        "domain": "cloudflare.com",
        "domainRegistered": "2009-02-17",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cloudflare.com/terms/",
        "privacy": "https://www.cloudflare.com/privacypolicy/",
        "statusPage": "https://www.cloudflarestatus.com",
        "changelog": "https://developers.cloudflare.com/ai-gateway/changelog/",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "notes": [
          "Terms are the Self-Serve Subscription Agreement, last updated 12 September 2025, which names Cloudflare, Inc. at 101 Townsend St, San Francisco.",
          "The service-specific terms (https://www.cloudflare.com/service-specific-terms-developer-platform/), last updated 28 September 2026, have a section for Workers AI and AI Gateway that supplements the agreement.",
          "The privacy policy is effective from 4 November 2025. The customer DPA is version 6.4, effective 3 April 2026.",
          "www.cloudflare.com/.well-known/security.txt, read on 9 October 2026, names HackerOne and a disclosure policy and has no Expires field. It is recorded as valid, although RFC 9116 asks for that field.",
          "The inference endpoints are on api.cloudflare.com and gateway.ai.cloudflare.com.",
          "The status page lists AI Gateway as its own component, operational on 9 October 2026.",
          "RDAP for cloudflare.com gives a registration date of 2009-02-17."
        ],
        "score": 95
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cloudflare-ai-gateway.json",
      "live": {
        "slug": "cloudflare-ai-gateway",
        "probe": {
          "target": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
          "method": "get",
          "lastAt": "2026-10-10T02:54:51.225095532Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 43,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 44,
          "p95ms24h": 109,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://www.cloudflarestatus.com",
          "indicator": "major",
          "summary": "Partial System Outage",
          "checkedAt": "2026-10-10T02:50:00.923581847Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "cloudflare/ai",
            "version": "ai-gateway-provider@4.0.1",
            "released": "2026-09-11",
            "seenAt": "2026-10-09T16:46:01.078937772Z"
          },
          {
            "registry": "npm",
            "name": "ai-gateway-provider",
            "version": "4.0.1",
            "seenAt": "2026-10-09T16:45:58.587383795Z"
          },
          {
            "registry": "npm",
            "name": "cloudflare",
            "version": "7.3.0",
            "seenAt": "2026-10-09T16:45:59.520764845Z"
          },
          {
            "registry": "pypi",
            "name": "cloudflare",
            "version": "5.9.0",
            "released": "2026-10-03",
            "seenAt": "2026-10-09T16:46:00.865337498Z"
          }
        ],
        "githubStars": 1189,
        "npmWeekly": 250648,
        "pypiWeekly": 365643,
        "pages": [
          {
            "url": "https://developers.cloudflare.com/ai-gateway/changelog/",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:34.168809494Z",
            "changedAt": "2026-10-07T18:04:13.791253735Z",
            "fingerprint": "aaad1064978c"
          },
          {
            "url": "https://developers.cloudflare.com/ai-gateway/reference/pricing/",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:35:36.234430797Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "57a018b683c5"
          },
          {
            "url": "https://www.cloudflare.com/privacypolicy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:49:06.00094733Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "e562d2a2da97"
          },
          {
            "url": "https://www.cloudflare.com/terms/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:49:08.260967052Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "d26dd2a74bde"
          }
        ],
        "updatedAt": "2026-10-10T02:54:51.225095532Z"
      }
    },
    "answer": "Cloudflare AI Gateway scores 66.8 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Prism Inference leads on schema \u0026 documentation.",
    "b": {
      "slug": "prism-inference",
      "name": "Prism Inference",
      "vendor": "Prism Technologies Inc",
      "vendorUrl": "https://prisminference.com",
      "kind": "model",
      "category": "inference",
      "summary": "Prism is a hosted inference API from Prism Technologies Inc for open-weight models, aimed at coding agents. It accepts OpenAI Chat Completions, OpenAI Responses and Anthropic Messages requests at api.prisminference.com. It launched on 24 September 2026.",
      "url": "https://www.anchorterminal.com/tools/prism-inference",
      "markdownUrl": "https://www.anchorterminal.com/tools/prism-inference.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/prism-inference.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/prism-inference.json",
      "repo": "https://github.com/prismhq/hermes-prism-provider",
      "license": "Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.prisminference.com/v1",
      "packages": [],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` with a key issued from account settings after sign-up at prisminference.com/signup. The inference endpoints also accept the key in `x-api-key`. A missing, invalid, expired or revoked key returns 401. No key scopes were found in the reviewed documentation. An agent can call `POST https://prisminference.com/api/agent-signups` with the owner's email and a username and receive a key once, but that key can't run inference until the owner supplies a six-digit emailed code, and it expires after 30 days. `GET /v1/models` needs no key.",
      "pricing": "usage",
      "pricingNotes": "Prepaid per-token pricing with no minimum. DeepSeek-V4.1-Flash is $0.09 in, $1.20 out and $0.06 cache read per million tokens, and Gemma 4 31B is $0.30, $0.40 and $0.15 (https://prisminference.com/pricing, matched by https://api.prisminference.com/v1/models). No free tier or trial credit was found, and a workspace without credit gets 402. A person funds the workspace in a browser. Elastic endpoints, dedicated deployments and batch are sold through sales with no published price.",
      "priceSummary": "from $0.09 / 1M in",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in llms.txt, the docs index, the OpenAPI file or the pricing page, read 2026-10-08. The docs describe prepaid credit funded by a person in the browser.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.prisminference.com",
      "rateLimitsUrl": "https://docs.prisminference.com/rate-limits",
      "llmsTxt": "https://prisminference.com/llms.txt",
      "openapi": "https://docs.prisminference.com/openapi.yaml",
      "capabilities": [
        "inference.fast",
        "inference.open-weights",
        "inference.llm"
      ],
      "tags": [
        "hosted",
        "model",
        "open-weights",
        "fast",
        "usage-priced",
        "prepaid",
        "openapi",
        "llms-txt",
        "openai-compatible",
        "anthropic-compatible",
        "zero-retention",
        "status-page",
        "new"
      ],
      "lastRelease": "2026-10-06",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 60.1,
        "grade": "C",
        "agentReady": false,
        "rank": 528,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 13,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 49,
          "payments": 30,
          "reliability": 65,
          "schema": 82,
          "security": 65,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": -2,
        "negativeNotes": [
          "2026-10-08. The home page shows a '99.99% Uptime SLA' tile, and the pricing page says 'No minimums, no rate limits' above the per-token table. The terms of 9 September 2026 say the services have no guaranteed uptime or service credit unless a separate written agreement says otherwise, and the docs describe per-key rate limits that return 429. No SLA document was found. A misleading claim, with the smallest deduction because the terms and docs state the real position (https://prisminference.com/, https://prisminference.com/pricing, https://prisminference.com/terms, https://docs.prisminference.com/rate-limits)."
        ],
        "verdict": "Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found.",
        "bestFor": "Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.",
        "strengths": [
          "One key works across OpenAI Chat Completions, OpenAI Responses and Anthropic Messages, with a public OpenAPI 3.1 file, llms.txt and Markdown docs",
          "Zero data retention is the default on every tier, and the privacy policy, terms and docs all say inputs and outputs are never used for training",
          "Every error carries a stable `code`, a `retryable` flag, a `fix` hint and a `docs_url`, and 429 carries `Retry-After` in seconds",
          "`GET /v1/models` answers without a key and returns context length, maximum output and per-token prices for each model",
          "DeepSeek-V4.1-Flash is listed with a 1M-token context and 384,000 output tokens at $0.09 in and $1.20 out per million"
        ],
        "weaknesses": [
          "Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available",
          "No rate-limit numbers are published. The docs say per-key limits exist, and the pricing page says 'no rate limits'",
          "The home page shows a '99.99% Uptime SLA' tile, while the terms say there is no guaranteed uptime without a separate written agreement",
          "No free tier found. Billing is prepaid, a person funds the workspace in a browser, and the agent sign-up key needs an emailed code",
          "The service launched on 24 September 2026. The changelog has one entry, and no deprecation policy, sub-processor list or certification was found"
        ],
        "agentNotes": [
          "Call `GET https://api.prisminference.com/v1/models` at start-up, with no key, and use only ids it returns. Expect 403 on `gemma-4-31b` without organisation access",
          "Use base URL `https://api.prisminference.com/v1` for OpenAI clients and `https://api.prisminference.com` with no `/v1` for Anthropic clients",
          "Read `error.retryable` before retrying, and wait for `Retry-After` on 429, which covers both key limits and model capacity",
          "Send `reasoning_effort: \"none\"` or `low` when latency matters. Reasoning is on by default and its tokens are billed as output",
          "Keep conversation state yourself and send `store: false` on Responses. `previous_response_id`, stored responses and hosted tools aren't supported"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 60.1
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 49,
          "payments": 30,
          "reliability": 65,
          "schema": 82,
          "security": 65,
          "transparency": 43
        },
        "provenanceScore": 79
      },
      "connect": {
        "install": "pip install openai   # or: npm install openai, base URL https://api.prisminference.com/v1. Anthropic SDKs use https://api.prisminference.com with no /v1",
        "http": "curl \"https://api.prisminference.com/v1/chat/completions\" \\\n  -H \"Authorization: Bearer $PRISM_API_KEY\" \\\n  -H \"Content-Type: application/json\" \\\n  -d '{\"model\":\"deepseek-v4.1-flash\",\"messages\":[{\"role\":\"user\",\"content\":\"Return pong.\"}]}'",
        "claudeCode": "export ANTHROPIC_BASE_URL=https://api.prisminference.com\nexport ANTHROPIC_AUTH_TOKEN=$PRISM_API_KEY\nexport ANTHROPIC_MODEL=deepseek-v4.1-flash\nexport ANTHROPIC_SMALL_FAST_MODEL=gemma-4-31b"
      },
      "letme": {
        "capability": "https://letme.dev/inference.fast",
        "tool": "https://letme.dev/prism-inference"
      },
      "area": "models",
      "provenance": {
        "legalEntity": "Prism Technologies Inc",
        "domain": "prisminference.com",
        "domainRegistered": "2026-09-09",
        "endpointOnVendorDomain": true,
        "terms": "https://prisminference.com/terms",
        "privacy": "https://prisminference.com/privacy",
        "statusPage": "https://status.prisminference.com",
        "changelog": "https://docs.prisminference.com/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "notes": [
          "The terms and the privacy policy, both last updated 9 September 2026, name Prism Technologies Inc. The terms are governed by California law with arbitration in San Francisco.",
          "RDAP gives a registration date of 2026-09-09 for prisminference.com, with Name.com as registrar.",
          "security.txt has a Contact line (founders@prisminference.com), an Expires date of 2027-10-06 and a Canonical line. No disclosure policy or bug bounty is named.",
          "The status page runs on incident.io with one component for each model and no component for the API or the website. Its incidents feed was empty on 8 October 2026.",
          "The YC directory lists Prism in the Spring 2025 batch, in San Francisco, with a team of 2. The home page links an X account named prism_videos, from the company's earlier video product."
        ],
        "score": 79
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/prism-inference.json",
      "live": {
        "slug": "prism-inference",
        "probe": {
          "target": "https://api.prisminference.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:06.706259541Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 131,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 149,
          "p95ms24h": 452,
          "samples24h": 249,
          "samples30d": 373,
          "days": [
            {
              "date": "2026-10-08",
              "probes": 93,
              "ok": 93
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.prisminference.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T02:50:43.671876128Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "prismhq/hermes-prism-provider",
            "version": "v1.0.2",
            "released": "2026-09-16",
            "seenAt": "2026-10-09T17:14:33.431257688Z"
          }
        ],
        "githubStars": 0,
        "securityTxt": {
          "url": "https://prisminference.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2027-10-06T00:00:00.000Z",
          "checkedAt": "2026-10-09T15:40:35.003140299Z"
        },
        "llmsTxt": {
          "url": "https://prisminference.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:37.937240222Z"
        },
        "pages": [
          {
            "url": "https://docs.prisminference.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-09T18:38:03.252133069Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "184a0fafdb8c"
          },
          {
            "url": "https://prisminference.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:52.542857787Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "dffb5d16a3c8"
          },
          {
            "url": "https://prisminference.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:54.77090078Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "64ea7b8f1b7b"
          },
          {
            "url": "https://prisminference.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:43:56.713185583Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "27dd15c69f49"
          }
        ],
        "updatedAt": "2026-10-10T02:55:06.706259541Z"
      }
    },
    "facts": [
      {
        "a": "Model router",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Cloudflare, Inc.",
        "b": "Prism Technologies Inc",
        "name": "Vendor"
      },
      {
        "a": "https://api.cloudflare.com/client/v4/accounts/{account_id}/ai",
        "b": "https://api.prisminference.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "payer tooling only",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under Cloudflare's Self-Serve Subscription Agreement and its service-specific terms. Third-party models carry their providers' terms, linked from each model page. `ai-gateway-provider` is MIT and Cloudflare's OpenAPI repository is BSD-3-Clause",
        "b": "Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-10-06",
        "b": "2026-10-06",
        "name": "Last release"
      },
      {
        "a": "2025-09-12",
        "b": "2026-09-09",
        "name": "Terms last updated"
      },
      {
        "a": "no date given",
        "b": "2026-09-09",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "251k npm/wk",
        "b": "none",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "Cloudflare AI Gateway scores 66.8 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Prism Inference leads on schema \u0026 documentation.",
        "question": "Which is better for AI agents, Cloudflare AI Gateway or Prism Inference?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Cloudflare AI Gateway and Prism Inference need an API key?"
      },
      {
        "answer": "Yes. Cloudflare AI Gateway has a hosted endpoint at https://api.cloudflare.com/client/v4/accounts/{account_id}/ai and Prism Inference at https://api.prisminference.com/v1.",
        "question": "Can an agent call Cloudflare AI Gateway and Prism Inference without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Agent ergonomics, 73 against 68",
          "Security \u0026 auth, 75 against 65",
          "Payments \u0026 pricing, 52 against 30",
          "Maintenance \u0026 community, 75 against 49",
          "Transparency \u0026 trust, 73 against 61"
        ],
        "also": null,
        "goodFor": "Teams that already hold a Cloudflare account and want one token, one bill, logs, budgets and fallbacks in front of several model providers, or that want to keep their own provider keys behind a proxy with spend limits.",
        "slug": "cloudflare-ai-gateway",
        "watchFor": "API tokens are account-scoped. Any token with AI Gateway Run can send requests through every gateway in the account and spend its stored provider keys"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 82 against 75"
        ],
        "also": null,
        "goodFor": "Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.",
        "slug": "prism-inference",
        "watchFor": "Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available"
      }
    ],
    "job": {
      "capability": "inference.llm",
      "name": "LLM inference"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-ai-gateway.json",
        "title": "Claude API vs Cloudflare AI Gateway",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-ai-gateway"
      },
      {
        "json": "https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference.json",
        "title": "Claude API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-prism-inference.json",
        "title": "Antseed vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-ai-gateway.json",
        "title": "BlockRun.AI vs Cloudflare AI Gateway",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-ai-gateway"
      },
      {
        "json": "https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference.json",
        "title": "BlockRun.AI vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai.json",
        "title": "Cloudflare AI Gateway vs Cloudflare Workers AI",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cohere-chat.json",
        "title": "Cloudflare AI Gateway vs Cohere Chat API (Command models)",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cohere-chat"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepinfra.json",
        "title": "Cloudflare AI Gateway vs DeepInfra",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepinfra"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepseek-api.json",
        "title": "Cloudflare AI Gateway vs DeepSeek API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepseek-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-gemini-api.json",
        "title": "Cloudflare AI Gateway vs Gemini Developer API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-gemini-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-groq.json",
        "title": "Cloudflare AI Gateway vs GroqCloud",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-groq"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-mistral-api.json",
        "title": "Cloudflare AI Gateway vs Mistral AI API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-mistral-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-novita-ai.json",
        "title": "Cloudflare AI Gateway vs Novita AI",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-novita-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openai-api.json",
        "title": "Cloudflare AI Gateway vs OpenAI API",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openai-api"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openrouter.json",
        "title": "Cloudflare AI Gateway vs OpenRouter",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openrouter"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-sambanova.json",
        "title": "Cloudflare AI Gateway vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-sambanova"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-siliconflow.json",
        "title": "Cloudflare AI Gateway vs SiliconFlow",
        "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-siliconflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.json",
        "title": "Cloudflare Workers AI vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference.json",
        "title": "Cohere Chat API (Command models) vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.json",
        "title": "DeepInfra vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference.json",
        "title": "DeepSeek API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference.json",
        "title": "Gemini Developer API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-vs-prism-inference.json",
        "title": "GroqCloud vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/groq-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference.json",
        "title": "Mistral AI API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference.json",
        "title": "Novita AI vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.json",
        "title": "OpenAI API vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/openai-api-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openrouter-vs-prism-inference.json",
        "title": "OpenRouter vs Prism Inference",
        "url": "https://www.anchorterminal.com/compare/openrouter-vs-prism-inference"
      },
      {
        "json": "https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow.json",
        "title": "Prism Inference vs SiliconFlow",
        "url": "https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/antseed-vs-cloudflare-ai-gateway.json",
        "title": "Antseed vs Cloudflare AI Gateway",
        "url": "https://www.anchorterminal.com/compare/antseed-vs-cloudflare-ai-gateway"
      },
      {
        "json": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.json",
        "title": "Prism Inference vs SambaCloud",
        "url": "https://www.anchorterminal.com/compare/prism-inference-vs-sambanova"
      }
    ],
    "scores": [
      {
        "by": 1,
        "cloudflare-ai-gateway": 66,
        "edge": "cloudflare-ai-gateway",
        "key": "reliability",
        "name": "Reliability",
        "prism-inference": 65,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "cloudflare-ai-gateway": 75,
        "edge": "prism-inference",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "prism-inference": 82,
        "weight": 13
      },
      {
        "by": 5,
        "cloudflare-ai-gateway": 73,
        "edge": "cloudflare-ai-gateway",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "prism-inference": 68,
        "weight": 13
      },
      {
        "by": 10,
        "cloudflare-ai-gateway": 75,
        "edge": "cloudflare-ai-gateway",
        "key": "security",
        "name": "Security \u0026 auth",
        "prism-inference": 65,
        "weight": 14
      },
      {
        "by": 22,
        "cloudflare-ai-gateway": 52,
        "edge": "cloudflare-ai-gateway",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "prism-inference": 30,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 26,
        "cloudflare-ai-gateway": 75,
        "edge": "cloudflare-ai-gateway",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "prism-inference": 49,
        "weight": 7
      },
      {
        "by": 12,
        "cloudflare-ai-gateway": 73,
        "edge": "cloudflare-ai-gateway",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "prism-inference": 61,
        "weight": 7
      }
    ],
    "summary": "Cloudflare AI Gateway scores 66.8 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Prism Inference leads on schema \u0026 documentation. Both do llm inference.",
    "verdicts": {
      "cloudflare-ai-gateway": "The gateway layer is free, per-model prices are public with no markup, and logs, spend limits and fallbacks are set per gateway. API tokens cannot be limited to one gateway, the service terms take third-party models outside Cloudflare's DPA, no SLA names the product, and error codes changed on 6 October 2026 with same-day notice.",
      "prism-inference": "Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference",
    "json": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.md",
    "slim": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.min.md"
  },
  "markdown": "Cloudflare AI Gateway scores 66.8 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Prism Inference leads on schema \u0026 documentation. Both do llm inference.\n\n- Cloudflare AI Gateway: grade B, 66.8/100, rank #275 of 950. Markdown https://www.anchorterminal.com/tools/cloudflare-ai-gateway.md · JSON https://www.anchorterminal.com/api/v1/tools/cloudflare-ai-gateway.json\n- Prism Inference: grade C, 60.1/100, rank #528 of 950. Markdown https://www.anchorterminal.com/tools/prism-inference.md · JSON https://www.anchorterminal.com/api/v1/tools/prism-inference.json\n- Best model APIs and inference for AI agents: https://www.anchorterminal.com/best/inference/index.md\n- All 136 models comparisons: https://www.anchorterminal.com/compare/inference/index.md\n\n## Which one, for what\n\n### Cloudflare AI Gateway (B)\n\nGood for: Teams that already hold a Cloudflare account and want one token, one bill, logs, budgets and fallbacks in front of several model providers, or that want to keep their own provider keys behind a proxy with spend limits.\n\nAhead on:\n- Agent ergonomics, 73 against 68\n- Security \u0026 auth, 75 against 65\n- Payments \u0026 pricing, 52 against 30\n- Maintenance \u0026 community, 75 against 49\n- Transparency \u0026 trust, 73 against 61\n\nWatch for: API tokens are account-scoped. Any token with AI Gateway Run can send requests through every gateway in the account and spend its stored provider keys\n\n### Prism Inference (C)\n\nGood for: Coding agents that want DeepSeek-V4.1-Flash at a low input price, with no retention, through whichever of the three wire formats the harness already speaks.\n\nAhead on:\n- Schema \u0026 documentation, 82 against 75\n\nWatch for: Two models. The docs mark Gemma 4 31B as request access per organisation, while llms.txt and the keyless catalogue list it as available\n\n\n## Score by category\n\n| Category | Weight | Cloudflare AI Gateway | Prism Inference | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 66 | 65 | Cloudflare AI Gateway +1 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 75 | 82 | Prism Inference +7 |\n| Agent ergonomics | 13% (16.2 this run) | 73 | 68 | Cloudflare AI Gateway +5 |\n| Security \u0026 auth | 14% (17.5 this run) | 75 | 65 | Cloudflare AI Gateway +10 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 52 | 30 | Cloudflare AI Gateway +22 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 75 | 49 | Cloudflare AI Gateway +26 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 73 | 61 | Cloudflare AI Gateway +12 |\n| Negative events | ≤15 | -3 | -2 | |\n| **Total** | | **66.8 · B** | **60.1 · C** | |\n\n## Facts side by side\n\n| Fact | Cloudflare AI Gateway | Prism Inference |\n| --- | --- | --- |\n| Kind | Model router | Model API |\n| Vendor | Cloudflare, Inc. | Prism Technologies Inc |\n| Hosted endpoint | `https://api.cloudflare.com/client/v4/accounts/{account_id}/ai` | `https://api.prisminference.com/v1` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | payer tooling only | no |\n| Licence | Proprietary service under Cloudflare's Self-Serve Subscription Agreement and its service-specific terms. Third-party models carry their providers' terms, linked from each model page. `ai-gateway-provider` is MIT and Cloudflare's OpenAPI repository is BSD-3-Clause | Proprietary service under Prism's terms of service. The OpenAPI file declares `LicenseRef-Proprietary`. The Hermes provider plugin repository carries no licence file |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-10-06 | 2026-10-06 |\n| Terms last updated | 2025-09-12 | 2026-09-09 |\n| Privacy policy last updated | no date given | 2026-09-09 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | yes | yes |\n| Terms restrict benchmarking | not found in the text | not found in the text |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | yes | yes |\n| Popularity | 251k npm/wk | none |\n\n## Verdicts\n\n**Cloudflare AI Gateway.** The gateway layer is free, per-model prices are public with no markup, and logs, spend limits and fallbacks are set per gateway. API tokens cannot be limited to one gateway, the service terms take third-party models outside Cloudflare's DPA, no SLA names the product, and error codes changed on 6 October 2026 with same-day notice.\n\n**Prism Inference.** Three wire formats, a public OpenAPI 3.1 file, per-token prices in a keyless catalogue and zero data retention by default on every tier. The service launched on 24 September 2026 with two models, one of them by request, from a two-person company. No rate-limit numbers, SLA document, free tier or deprecation policy was found.\n\n## Before you call either\n\n### Cloudflare AI Gateway\n\n1. Use the REST API at `api.cloudflare.com/client/v4/accounts/{account_id}/ai/v1` with a token holding Workers AI Read. A token with only an AI Gateway permission returns 401 with code 10000\n2. Send `cf-aig-gateway-id` on every Workers AI (`@cf/`) call and on dynamic routes. Without it a dynamic route resolves against the default gateway and returns 404\n3. Set `cf-aig-collect-log: false` or `cf-aig-collect-log-payload: false` on sensitive requests. Logging of prompts and responses is on by default\n4. Treat a 429 as a rate limit or a spent budget and a 401 with code 2009 as a rejected provider key. Do not retry a 2009\n5. Set `byok_only` on the gateway or send `cf-aig-no-wholesale: true` to stop a request without a provider key falling through to paid Unified Billing credits\n\n### Prism Inference\n\n1. Call `GET https://api.prisminference.com/v1/models` at start-up, with no key, and use only ids it returns. Expect 403 on `gemma-4-31b` without organisation access\n2. Use base URL `https://api.prisminference.com/v1` for OpenAI clients and `https://api.prisminference.com` with no `/v1` for Anthropic clients\n3. Read `error.retryable` before retrying, and wait for `Retry-After` on 429, which covers both key limits and model capacity\n4. Send `reasoning_effort: \"none\"` or `low` when latency matters. Reasoning is on by default and its tokens are billed as output\n5. Keep conversation state yourself and send `store: false` on Responses. `previous_response_id`, stored responses and hosted tools aren't supported\n\n## Questions\n\n### Which is better for AI agents, Cloudflare AI Gateway or Prism Inference?\n\nCloudflare AI Gateway scores 66.8 (B) on agent readiness against Prism Inference's 60.1 (C), and leads in 6 of 7 scored categories. Prism Inference leads on schema \u0026 documentation.\n\n### Do Cloudflare AI Gateway and Prism Inference need an API key?\n\nBoth need an API key.\n\n### Can an agent call Cloudflare AI Gateway and Prism Inference without installing anything?\n\nYes. Cloudflare AI Gateway has a hosted endpoint at https://api.cloudflare.com/client/v4/accounts/{account_id}/ai and Prism Inference at https://api.prisminference.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cloudflare-ai-gateway\", \"b\": \"prism-inference\"}`. From a terminal: `anchor compare cloudflare-ai-gateway prism-inference`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cloudflare-ai-gateway.json and https://www.anchorterminal.com/api/v1/tools/prism-inference.json\n\n## Other comparisons with Cloudflare AI Gateway or Prism Inference\n\n- [Claude API vs Cloudflare AI Gateway](https://www.anchorterminal.com/compare/anthropic-api-vs-cloudflare-ai-gateway.md)\n- [Claude API vs Prism Inference](https://www.anchorterminal.com/compare/anthropic-api-vs-prism-inference.md)\n- [Antseed vs Prism Inference](https://www.anchorterminal.com/compare/antseed-vs-prism-inference.md)\n- [BlockRun.AI vs Cloudflare AI Gateway](https://www.anchorterminal.com/compare/blockrun-ai-vs-cloudflare-ai-gateway.md)\n- [BlockRun.AI vs Prism Inference](https://www.anchorterminal.com/compare/blockrun-ai-vs-prism-inference.md)\n- [Cloudflare AI Gateway vs Cloudflare Workers AI](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cloudflare-workers-ai.md)\n- [Cloudflare AI Gateway vs Cohere Chat API (Command models)](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-cohere-chat.md)\n- [Cloudflare AI Gateway vs DeepInfra](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepinfra.md)\n- [Cloudflare AI Gateway vs DeepSeek API](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-deepseek-api.md)\n- [Cloudflare AI Gateway vs Gemini Developer API](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-gemini-api.md)\n- [Cloudflare AI Gateway vs GroqCloud](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-groq.md)\n- [Cloudflare AI Gateway vs Mistral AI API](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-mistral-api.md)\n- [Cloudflare AI Gateway vs Novita AI](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-novita-ai.md)\n- [Cloudflare AI Gateway vs OpenAI API](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openai-api.md)\n- [Cloudflare AI Gateway vs OpenRouter](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-openrouter.md)\n- [Cloudflare AI Gateway vs SambaCloud](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-sambanova.md)\n- [Cloudflare AI Gateway vs SiliconFlow](https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-siliconflow.md)\n- [Cloudflare Workers AI vs Prism Inference](https://www.anchorterminal.com/compare/cloudflare-workers-ai-vs-prism-inference.md)\n- [Cohere Chat API (Command models) vs Prism Inference](https://www.anchorterminal.com/compare/cohere-chat-vs-prism-inference.md)\n- [DeepInfra vs Prism Inference](https://www.anchorterminal.com/compare/deepinfra-vs-prism-inference.md)\n- [DeepSeek API vs Prism Inference](https://www.anchorterminal.com/compare/deepseek-api-vs-prism-inference.md)\n- [Gemini Developer API vs Prism Inference](https://www.anchorterminal.com/compare/gemini-api-vs-prism-inference.md)\n- [GroqCloud vs Prism Inference](https://www.anchorterminal.com/compare/groq-vs-prism-inference.md)\n- [Mistral AI API vs Prism Inference](https://www.anchorterminal.com/compare/mistral-api-vs-prism-inference.md)\n- [Novita AI vs Prism Inference](https://www.anchorterminal.com/compare/novita-ai-vs-prism-inference.md)\n- [OpenAI API vs Prism Inference](https://www.anchorterminal.com/compare/openai-api-vs-prism-inference.md)\n- [OpenRouter vs Prism Inference](https://www.anchorterminal.com/compare/openrouter-vs-prism-inference.md)\n- [Prism Inference vs SiliconFlow](https://www.anchorterminal.com/compare/prism-inference-vs-siliconflow.md)\n- [Antseed vs Cloudflare AI Gateway](https://www.anchorterminal.com/compare/antseed-vs-cloudflare-ai-gateway.md)\n- [Prism Inference vs SambaCloud](https://www.anchorterminal.com/compare/prism-inference-vs-sambanova.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cloudflare AI Gateway vs Prism Inference",
        "url": ""
      }
    ],
    "description": "Cloudflare AI Gateway scores 66.8 (B) to Prism Inference's 60.1 (C) for llm inference. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Cloudflare AI Gateway B 66.8",
      "Prism Inference C 60.1",
      "scores"
    ],
    "h1": "Cloudflare AI Gateway vs Prism Inference",
    "image": "https://www.anchorterminal.com/assets/og/compare-cloudflare-ai-gateway-vs-prism-inference.png",
    "path": "/compare/cloudflare-ai-gateway-vs-prism-inference",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cloudflare AI Gateway vs Prism Inference for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/cloudflare-ai-gateway-vs-prism-inference"
  },
  "tokens": {
    "markdown": 3000,
    "slim": 730
  },
  "version": 1
}
