{
  "data": {
    "a": {
      "slug": "nebius-ai-cloud",
      "name": "Nebius AI Cloud",
      "vendor": "Nebius",
      "vendorUrl": "https://nebius.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Nebius AI Cloud rents NVIDIA GPU virtual machines and InfiniBand clusters, with managed Kubernetes, Slurm and Serverless AI jobs and endpoints for containers. Resources are managed through REST and gRPC APIs, a CLI, a Terraform provider and SDKs.",
      "url": "https://www.anchorterminal.com/tools/nebius-ai-cloud",
      "markdownUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json",
      "repo": "https://github.com/nebius/api",
      "license": "Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT",
      "transports": [
        "http",
        "stdio"
      ],
      "remoteUrl": "https://api.nebius.cloud",
      "packages": [
        {
          "registry": "pypi",
          "name": "nebius"
        },
        {
          "registry": "npm",
          "name": "@nebius/js-sdk"
        },
        {
          "registry": "go",
          "name": "github.com/nebius/gosdk"
        }
      ],
      "auth": "mixed",
      "authNotes": "Self-serve. A person signs up in the web console with a Google, GitHub or Microsoft account. Every API call takes `Authorization: Bearer` with an access token valid for 12 hours. A user gets one from `nebius iam get-access-token`. A service account uploads an RSA public key (an authorised key, with optional expiry), signs a five-minute RS256 JWT and exchanges it at `https://auth.eu.nebius.com/oauth2/token/exchange`. Permissions come from group roles (`auditor`, `viewer`, `editor`, `admin` and service roles) granted on a tenant, project or resource. Serverless AI endpoints take their own token set in `spec.authToken`.",
      "pricing": "usage",
      "pricingNotes": "Pay as you go, billed by the second, with no free tier or trial found. On-demand per GPU-hour from 1 October 2026, B300 $9.50, B200 $8.50, H200 $5.40, H100 $4.50, RTX PRO 6000 $1.80, and L40S $1.35 plus vCPU and RAM. Preemptible GPUs are spot priced from $0.79. Serverless AI bills at Compute prices and a stopped endpoint bills nothing. Adding a card charges $25 to the balance. Commitment discounts go through sales (https://docs.nebius.com/compute/resources/pricing, https://nebius.com/prices).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the OpenAPI document or the price list (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 4137,
        "pypiWeekly": 468852,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://docs.nebius.com/",
      "llmsTxt": "https://docs.nebius.com/llms.txt",
      "openapi": "https://api.nebius.cloud/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.batch",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "openapi",
        "llms-txt",
        "grpc",
        "terraform",
        "mcp",
        "python",
        "typescript",
        "go",
        "status-page",
        "soc2",
        "enterprise"
      ],
      "lastRelease": "2026-10-07",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67.2,
        "grade": "B",
        "agentReady": false,
        "rank": 240,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 1,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 77,
          "maintenance": 82,
          "payments": 20,
          "reliability": 57,
          "schema": 78,
          "security": 81,
          "transparency": 77
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card.",
        "bestFor": "Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.",
        "strengths": [
          "OpenAPI 3.0.3 document at `https://api.nebius.cloud/openapi.json` with 602 operations, generated from the same protobuf definitions as the gRPC API, CLI, Terraform provider and SDKs",
          "`X-Idempotency-Key` header for modifying calls, and a `retry_type` field on errors that says whether to retry the call",
          "Service accounts sign in with an uploaded RSA key and receive 12-hour tokens, with roles granted per tenant, project or resource",
          "Per-second billing with public prices, and a 99.5 per cent monthly uptime commitment per virtual machine",
          "llms.txt, every docs page as Markdown, and a keyless docs MCP server at `https://docs.nebius.com/mcp`"
        ],
        "weaknesses": [
          "Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August",
          "No request rate limits with numbers and no Retry-After guidance found in the reviewed documentation",
          "Serverless AI endpoints run on one container VM that is started and stopped by hand; no autoscaling or scale to zero found",
          "No free tier or trial found. Adding a card at signup charges $25 to the balance, and signup is a browser flow through Google, GitHub or Microsoft",
          "The OpenAPI document has no examples, documents only 200 responses and reports its version as `version not set`"
        ],
        "agentNotes": [
          "Use a service account with an authorised key, then exchange a five-minute RS256 JWT at `https://auth.eu.nebius.com/oauth2/token/exchange` for a 12-hour Bearer token",
          "Send `X-Idempotency-Key` with a random UUID on every create, update and delete, since a 504 can follow a call that succeeded",
          "Poll the returned operation (`/ai/v1/endpoints/operations/{id}`) until `status` is set; concurrent operations on one resource are not supported",
          "Stop or delete endpoints when idle. A stopped endpoint bills nothing, a stopped Devlab or VM still bills for its disk",
          "Check region support first. Serverless AI is absent from `eu-south1` and `us-north1`, and each GPU platform exists in one to four regions",
          "Run the beta `nebius/mcp-server` with safe mode on (the default); `nebius_cli_execute` can run any CLI command when `SAFE_MODE=false`"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67.2
          }
        ],
        "editorialScores": {
          "ergonomics": 77,
          "maintenance": 82,
          "payments": 20,
          "reliability": 57,
          "schema": 78,
          "security": 81,
          "transparency": 70
        },
        "provenanceScore": 83
      },
      "connect": {
        "install": "curl -sSL https://artifacts.nebius.cloud/cli/install.sh | bash",
        "http": "curl --request GET --url 'https://api.nebius.cloud/iam/v1/profiles' --header 'Authorization: Bearer \u003caccess_token\u003e'",
        "config": {
          "mcpServers": {
            "Nebius MCP Server": {
              "args": [
                "--refresh-package",
                "nebius-mcp-server",
                "nebius-mcp-server@git+https://github.com/nebius/mcp-server@main"
              ],
              "command": "uvx",
              "env": {}
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/nebius-ai-cloud"
      },
      "sameCompany": [
        "nebius-token-factory-fine-tuning"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "NVIDIA H200 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 5.4,
          "note": "Billed per second. $4.50 before 1 October 2026"
        },
        {
          "item": "NVIDIA H100 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 4.5,
          "note": "$3.85 before 1 October 2026"
        },
        {
          "item": "NVIDIA B200 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 8.5
        },
        {
          "item": "NVIDIA B300 NVLink, on demand",
          "unit": "gpu-hour",
          "usd": 9.5
        },
        {
          "item": "NVIDIA RTX PRO 6000, on demand",
          "unit": "gpu-hour",
          "usd": 1.8
        },
        {
          "item": "NVIDIA L40S, GPU only",
          "unit": "gpu-hour",
          "usd": 1.35,
          "note": "vCPU ($0.01 to $0.012 an hour) and RAM ($0.0032 a GiB-hour) are billed separately"
        },
        {
          "item": "NVIDIA H200 NVLink, preemptible",
          "unit": "gpu-hour",
          "usd": 0.79,
          "note": "Minimum spot price from 8 October 2026. The spot price can change every 15 minutes"
        },
        {
          "item": "Network SSD disk",
          "unit": "gb-month",
          "usd": 0.071,
          "note": "Per GiB for 730 hours"
        }
      ],
      "provenance": {
        "legalEntity": "Nebius B.V.",
        "domain": "nebius.com",
        "domainRegistered": "2004-06-26",
        "endpointOnVendorDomain": false,
        "terms": "https://docs.nebius.com/legal/agreement",
        "privacy": "https://docs.nebius.com/legal/privacy",
        "statusPage": "https://status.nebius.com",
        "changelog": "https://docs.nebius.com/cli/release-notes",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "notes": [
          "The Services Agreement (published 15 September 2026, effective 28 September 2026) names Nebius B.V. under Dutch law for customers outside the United States and Israel, with other Nebius entities for those two countries. The separate Terms of Use page covers the website.",
          "The privacy policy (23 September 2026) gives Nebius B.V., Burgerweeshuispad 101, 1076ER Amsterdam, and says data processed for customers as a processor falls under the DPA at https://docs.nebius.com/legal/dpa.",
          "The API answers at api.nebius.cloud and tokens are exchanged at auth.eu.nebius.com. nebius.cloud is a second domain that Nebius's docs name for the API and the CLI installer.",
          "security.txt at nebius.com gives security@nebius.com and expires 2027-12-31. It has no Policy field.",
          "The changelog link is the CLI release notes, which are generated from the API. No separate API changelog was found.",
          "RDAP for nebius.com gives a registration date of 2004-06-26."
        ],
        "score": 83
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/nebius-ai-cloud.json",
      "live": {
        "slug": "nebius-ai-cloud",
        "probe": {
          "target": "https://api.nebius.cloud",
          "method": "get",
          "lastAt": "2026-10-09T09:26:58.348144222Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 246,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 158,
          "p95ms24h": 204,
          "samples24h": 20,
          "samples30d": 20,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 20,
              "ok": 20
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.nebius.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T09:25:19.246706859Z"
        },
        "updatedAt": "2026-10-09T09:26:58.348144222Z"
      }
    },
    "answer": "Nebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing.",
    "b": {
      "slug": "replicate-deploy",
      "name": "Replicate Deployments",
      "vendor": "Replicate",
      "vendorUrl": "https://replicate.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Replicate's service for deploying and running custom models.",
      "url": "https://www.anchorterminal.com/tools/replicate-deploy",
      "markdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json",
      "repo": "https://github.com/replicate/cog",
      "license": "Apache-2.0",
      "transports": [
        "http",
        "sse",
        "stdio"
      ],
      "remoteUrl": "https://api.replicate.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "replicate"
        },
        {
          "registry": "pypi",
          "name": "replicate"
        },
        {
          "registry": "npm",
          "name": "replicate-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API token on every call to api.replicate.com. `cog push` uses the same token to upload a model image. The hosted MCP at https://mcp.replicate.com/sse asks for the token in a browser flow and holds it for the client; the local `replicate-mcp` package reads `REPLICATE_API_TOKEN`.",
      "pricing": "usage",
      "pricingNotes": "Private models and deployments bill per second for the whole time an instance is up, set-up and idle included, from prepaid credit or monthly in arrears. CPU $0.000100 a second ($0.36 an hour), T4 $0.000225 ($0.81), L40S $0.000975 ($3.51), A100 80 GB $0.001400 ($5.04), H100 $0.001525 ($5.49), 2x L40S $0.001950 ($7.02), 2x A100 $0.002800 ($10.08). 2x H100 ($10.98), 4x and 8x L40S, A100 and H100 up to $43.92 an hour need a committed-spend contract. Fast-booting fine-tunes bill only while active. Public models bill only active time and not failures (https://replicate.com/pricing, https://replicate.com/docs/topics/billing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 9500,
        "npmWeekly": 634116,
        "pypiWeekly": 386704,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://replicate.com/docs/topics/deployments",
      "llmsTxt": "https://replicate.com/docs/llms.txt",
      "openapi": "https://api.replicate.com/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "mcp",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "async-jobs",
        "webhooks",
        "open-source"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63.6,
        "grade": "B",
        "agentReady": false,
        "rank": 347,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 5,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 78
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.",
        "bestFor": "Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.",
        "strengths": [
          "OpenAPI file, llms.txt and an MCP server with a two-tool code mode",
          "Deployment min and max instances settable over the API, 0 allowed",
          "API prediction data deleted after one hour by default",
          "Published limits, 600 prediction creates and 3,000 other calls a minute",
          "Leaked tokens found on GitHub are disabled automatically"
        ],
        "weaknesses": [
          "Private instances bill set-up and idle time, H100 at $5.49 an hour",
          "API tokens have no scopes, expiry or audit log",
          "Changelog silent since 21 April 2026",
          "Only T4, L40S, A100 and H100, and more than 2 GPUs needs a committed-spend contract",
          "Two September 2026 incidents ran 15 and 20 hours, both marked minor"
        ],
        "agentNotes": [
          "List `GET /v1/hardware` first and use the returned `sku` in the deployment body",
          "Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not",
          "Send `Prefer: wait` on deployment predictions to block instead of polling",
          "Copy outputs within an hour; API prediction data is deleted after that",
          "Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63.6
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 69
        },
        "provenanceScore": 87
      },
      "connect": {
        "install": "pip install cog replicate",
        "http": "curl -X POST \"https://api.replicate.com/v1/deployments/$REPLICATE_OWNER/my-deployment/predictions\" \\\n  -H \"Authorization: Bearer $REPLICATE_API_TOKEN\" -H \"Content-Type: application/json\" -H \"Prefer: wait\" \\\n  -d '{\"input\":{\"prompt\":\"hello\"}}'",
        "claudeCode": "claude mcp add replicate https://mcp.replicate.com/sse --transport sse --scope user",
        "config": {
          "mcpServers": {
            "replicate": {
              "args": [
                "-y",
                "replicate-mcp"
              ],
              "command": "npx",
              "env": {
                "REPLICATE_API_TOKEN": "${REPLICATE_API_TOKEN}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/replicate-deploy"
      },
      "sameCompany": [
        "replicate-image",
        "replicate-video",
        "replicate-musicgen"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.49,
          "note": "$0.001525 a second, including set-up and idle"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.04,
          "note": "$0.001400 a second"
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 3.51,
          "note": "$0.000975 a second"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.81,
          "note": "$0.000225 a second"
        }
      ],
      "provenance": {
        "legalEntity": "Replicate, LLC",
        "domain": "replicate.com",
        "domainRegistered": "1998-05-26",
        "domainNote": "replicate.com was registered in 1998, long before Replicate the company existed.",
        "endpointOnVendorDomain": true,
        "terms": "https://replicate.com/terms",
        "privacy": "https://replicate.com/privacy",
        "statusPage": "https://replicatestatus.com",
        "changelog": "https://replicate.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms last updated 2026-04-01 name Replicate, LLC as the contracting party.",
          "replicatestatus.com redirects to Cloudflare's status page filtered to Replicate.",
          "Replicate's hosted image and music models are listed separately under image generation and music generation."
        ],
        "score": 87
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/replicate-deploy.json",
      "live": {
        "slug": "replicate-deploy",
        "probe": {
          "target": "https://api.replicate.com/v1",
          "method": "get",
          "lastAt": "2026-10-09T09:27:02.672787909Z",
          "lastOk": true,
          "lastStatus": 401,
          "lastMs": 143,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 150,
          "p95ms24h": 327,
          "samples24h": 261,
          "samples30d": 2085,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 101,
              "ok": 101
            }
          ]
        },
        "vendorStatus": {
          "page": "https://replicatestatus.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-09T07:58:28.982306027Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "replicate/cog",
            "version": "v0.23.0",
            "released": "2026-09-22",
            "seenAt": "2026-10-08T16:27:09.246140372Z"
          },
          {
            "registry": "npm",
            "name": "replicate",
            "version": "1.4.0",
            "seenAt": "2026-10-08T16:27:06.602404036Z"
          },
          {
            "registry": "npm",
            "name": "replicate-mcp",
            "version": "0.9.0",
            "seenAt": "2026-10-08T16:27:09.031273281Z"
          },
          {
            "registry": "pypi",
            "name": "replicate",
            "version": "1.0.7",
            "released": "2025-05-27",
            "seenAt": "2026-10-08T16:27:07.620263226Z"
          }
        ],
        "githubStars": 9487,
        "npmWeekly": 704939,
        "pypiWeekly": 356876,
        "securityTxt": {
          "url": "https://replicate.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:39:00.17768337Z"
        },
        "llmsTxt": {
          "url": "https://replicate.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:48.463855942Z"
        },
        "domain": {
          "domain": "replicate.com",
          "registered": "1998-05-26",
          "source": "https://rdap.verisign.com/com/v1/domain/replicate.com",
          "checkedAt": "2026-10-04T13:07:04.742407865Z"
        },
        "pages": [
          {
            "url": "https://replicate.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-08T18:23:40.627013731Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "490f4836aca3"
          },
          {
            "url": "https://replicate.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:43.001722065Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3f1305f154be"
          },
          {
            "url": "https://replicate.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:45.469231625Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "8e299fbc64eb"
          },
          {
            "url": "https://replicate.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-08T18:23:46.848375869Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ea48efe3382b"
          }
        ],
        "updatedAt": "2026-10-09T09:27:02.672787909Z"
      }
    },
    "facts": [
      {
        "a": "HTTP API",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Nebius",
        "b": "Replicate",
        "name": "Vendor"
      },
      {
        "a": "https://api.nebius.cloud",
        "b": "https://api.replicate.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, stdio",
        "b": "HTTP, SSE (legacy), stdio",
        "name": "Transports"
      },
      {
        "a": "OAuth or key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "$1.35 per GPU-hour",
        "b": "not published",
        "name": "Price for compute gpu"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT",
        "b": "Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-10-07",
        "b": "2026-09-22",
        "name": "Last release"
      },
      {
        "a": "2026-09-28",
        "b": "2026-04-01",
        "name": "Terms last updated"
      },
      {
        "a": "2026-09-23",
        "b": "2026-04-01",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "4.1k npm/wk, 469k PyPI/wk",
        "b": "9.5k stars, 634k npm/wk, 387k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Nebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing.",
        "question": "Which is better for AI agents, Nebius AI Cloud or Replicate Deployments?"
      },
      {
        "answer": "Nebius AI Cloud takes an API key or an OAuth sign-in. Replicate Deployments needs an API key.",
        "question": "Do Nebius AI Cloud and Replicate Deployments need an API key?"
      },
      {
        "answer": "Yes. Nebius AI Cloud has a hosted endpoint at https://api.nebius.cloud and Replicate Deployments at https://api.replicate.com/v1.",
        "question": "Can an agent call Nebius AI Cloud and Replicate Deployments without installing anything?"
      },
      {
        "answer": "No open-source release is listed for Nebius AI Cloud. Replicate Deployments is open source (Apache-2.0).",
        "question": "Are Nebius AI Cloud and Replicate Deployments open source?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Agent ergonomics, 77 against 68",
          "Security \u0026 auth, 81 against 40",
          "Maintenance \u0026 community, 82 against 70"
        ],
        "also": null,
        "goodFor": "Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.",
        "slug": "nebius-ai-cloud",
        "watchFor": "Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August"
      },
      {
        "aheadOn": [
          "Reliability, 75 against 57",
          "Schema \u0026 documentation, 85 against 78",
          "Payments \u0026 pricing, 30 against 20"
        ],
        "also": [
          "Open source"
        ],
        "goodFor": "Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.",
        "slug": "replicate-deploy",
        "watchFor": "Private instances bill set-up and idle time, H100 at $5.49 an hour"
      }
    ],
    "job": {
      "capability": "compute.gpu",
      "name": "Compute gpu"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.json",
        "title": "Baseten vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.json",
        "title": "Baseten vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud.json",
        "title": "Beam vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.json",
        "title": "Beam vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/beam-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud.json",
        "title": "Cerebrium vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.json",
        "title": "Cerebrium vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud.json",
        "title": "CoreWeave vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy.json",
        "title": "CoreWeave vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud.json",
        "title": "Hugging Face Inference Endpoints vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-replicate-deploy.json",
        "title": "Hugging Face Inference Endpoints vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud.json",
        "title": "Hyperbolic vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/hyperbolic-vs-replicate-deploy.json",
        "title": "Hyperbolic vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/hyperbolic-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud.json",
        "title": "Koyeb vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.json",
        "title": "Koyeb vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud.json",
        "title": "Lambda Cloud vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.json",
        "title": "Lambda Cloud vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud.json",
        "title": "Modal vs Nebius AI Cloud",
        "url": "https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud"
      },
      {
        "json": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.json",
        "title": "Modal vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/modal-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank.json",
        "title": "Nebius AI Cloud vs Northflank",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod.json",
        "title": "Nebius AI Cloud vs Runpod",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute.json",
        "title": "Nebius AI Cloud vs Thunder Compute",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai.json",
        "title": "Nebius AI Cloud vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda.json",
        "title": "Nebius AI Cloud vs Verda",
        "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda"
      },
      {
        "json": "https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.json",
        "title": "Northflank vs Replicate Deployments",
        "url": "https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.json",
        "title": "Replicate Deployments vs Runpod",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-thunder-compute.json",
        "title": "Replicate Deployments vs Thunder Compute",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-thunder-compute"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.json",
        "title": "Replicate Deployments vs Vast.ai",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/replicate-deploy-vs-verda.json",
        "title": "Replicate Deployments vs Verda",
        "url": "https://www.anchorterminal.com/compare/replicate-deploy-vs-verda"
      }
    ],
    "scores": [
      {
        "by": 18,
        "edge": "replicate-deploy",
        "key": "reliability",
        "name": "Reliability",
        "nebius-ai-cloud": 57,
        "replicate-deploy": 75,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "edge": "replicate-deploy",
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "nebius-ai-cloud": 78,
        "replicate-deploy": 85,
        "weight": 13
      },
      {
        "by": 9,
        "edge": "nebius-ai-cloud",
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "nebius-ai-cloud": 77,
        "replicate-deploy": 68,
        "weight": 13
      },
      {
        "by": 41,
        "edge": "nebius-ai-cloud",
        "key": "security",
        "name": "Security \u0026 auth",
        "nebius-ai-cloud": 81,
        "replicate-deploy": 40,
        "weight": 14
      },
      {
        "by": 10,
        "edge": "replicate-deploy",
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "nebius-ai-cloud": 20,
        "replicate-deploy": 30,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 12,
        "edge": "nebius-ai-cloud",
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "nebius-ai-cloud": 82,
        "replicate-deploy": 70,
        "weight": 7
      },
      {
        "by": 1,
        "edge": "replicate-deploy",
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "nebius-ai-cloud": 77,
        "replicate-deploy": 78,
        "weight": 7
      }
    ],
    "summary": "Nebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing. Both do compute gpu.",
    "verdicts": {
      "nebius-ai-cloud": "One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card.",
      "replicate-deploy": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy",
    "json": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.md",
    "slim": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.min.md"
  },
  "markdown": "Nebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing. Both do compute gpu.\n\n- Nebius AI Cloud: grade B, 67.2/100, rank #240 of 842. Markdown https://www.anchorterminal.com/tools/nebius-ai-cloud.md · JSON https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json\n- Replicate Deployments: grade B, 63.6/100, rank #347 of 842. Markdown https://www.anchorterminal.com/tools/replicate-deploy.md · JSON https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Which one, for what\n\n### Nebius AI Cloud (B)\n\nGood for: Teams that want whole GPU VMs or InfiniBand clusters in Europe, the UK, Israel or the US with IAM, Terraform and an SLA, and are content to manage endpoint lifecycles themselves.\n\nAhead on:\n- Agent ergonomics, 77 against 68\n- Security \u0026 auth, 81 against 40\n- Maintenance \u0026 community, 82 against 70\n\nWatch for: Status page lists 14 incidents marked major between 14 July and 8 October 2026, including about 21 hours of partial degradation in us-central1 on 19 August\n\n### Replicate Deployments (B)\n\nGood for: Teams already calling Replicate's public models who want their own model behind the same API, MCP server and webhooks.\n\nAhead on:\n- Reliability, 75 against 57\n- Schema \u0026 documentation, 85 against 78\n- Payments \u0026 pricing, 30 against 20\n\nAlso in its favour:\n- Open source\n\nWatch for: Private instances bill set-up and idle time, H100 at $5.49 an hour\n\n\n## Score by category\n\n| Category | Weight | Nebius AI Cloud | Replicate Deployments | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 57 | 75 | Replicate Deployments +18 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 78 | 85 | Replicate Deployments +7 |\n| Agent ergonomics | 13% (16.2 this run) | 77 | 68 | Nebius AI Cloud +9 |\n| Security \u0026 auth | 14% (17.5 this run) | 81 | 40 | Nebius AI Cloud +41 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 30 | Replicate Deployments +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 82 | 70 | Nebius AI Cloud +12 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 77 | 78 | Replicate Deployments +1 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **67.2 · B** | **63.6 · B** | |\n\n## Facts side by side\n\n| Fact | Nebius AI Cloud | Replicate Deployments |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Nebius | Replicate |\n| Hosted endpoint | `https://api.nebius.cloud` | `https://api.replicate.com/v1` |\n| Transports | HTTP, stdio | HTTP, SSE (legacy), stdio |\n| Auth | OAuth or key | API key |\n| Pricing | Pay per use | Pay per use |\n| Price for compute gpu | $1.35 per GPU-hour | not published |\n| x402 | no | no |\n| Licence | Proprietary service under the Nebius Services Agreement. The API definitions, the Go, Python and JavaScript SDKs and the MCP server on GitHub are MIT | Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-10-07 | 2026-09-22 |\n| Terms last updated | 2026-09-28 | 2026-04-01 |\n| Privacy policy last updated | 2026-09-23 | 2026-04-01 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | yes | not found in the text |\n| Terms or service can change without notice | not found in the text | yes |\n| Arbitration or class-action waiver | yes | yes |\n| Popularity | 4.1k npm/wk, 469k PyPI/wk | 9.5k stars, 634k npm/wk, 387k PyPI/wk |\n| Agent reviews | none | 3/5 (2) |\n\n## Verdicts\n\n**Nebius AI Cloud.** One API definition generates the REST and gRPC interfaces, the CLI, Terraform provider and three SDKs, with a 602-operation OpenAPI document, `X-Idempotency-Key` and role-scoped service accounts. The status page lists 14 major incidents between 14 July and 8 October 2026, no request rate limits were found, and signup needs a browser and a card.\n\n**Replicate Deployments.** OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.\n\n## Before you call either\n\n### Nebius AI Cloud\n\n1. Use a service account with an authorised key, then exchange a five-minute RS256 JWT at `https://auth.eu.nebius.com/oauth2/token/exchange` for a 12-hour Bearer token\n2. Send `X-Idempotency-Key` with a random UUID on every create, update and delete, since a 504 can follow a call that succeeded\n3. Poll the returned operation (`/ai/v1/endpoints/operations/{id}`) until `status` is set; concurrent operations on one resource are not supported\n4. Stop or delete endpoints when idle. A stopped endpoint bills nothing, a stopped Devlab or VM still bills for its disk\n5. Check region support first. Serverless AI is absent from `eu-south1` and `us-north1`, and each GPU platform exists in one to four regions\n6. Run the beta `nebius/mcp-server` with safe mode on (the default); `nebius_cli_execute` can run any CLI command when `SAFE_MODE=false`\n\n### Replicate Deployments\n\n1. List `GET /v1/hardware` first and use the returned `sku` in the deployment body\n2. Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not\n3. Send `Prefer: wait` on deployment predictions to block instead of polling\n4. Copy outputs within an hour; API prediction data is deleted after that\n5. Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute\n\n## Questions\n\n### Which is better for AI agents, Nebius AI Cloud or Replicate Deployments?\n\nNebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing.\n\n### Do Nebius AI Cloud and Replicate Deployments need an API key?\n\nNebius AI Cloud takes an API key or an OAuth sign-in. Replicate Deployments needs an API key.\n\n### Can an agent call Nebius AI Cloud and Replicate Deployments without installing anything?\n\nYes. Nebius AI Cloud has a hosted endpoint at https://api.nebius.cloud and Replicate Deployments at https://api.replicate.com/v1.\n\n### Are Nebius AI Cloud and Replicate Deployments open source?\n\nNo open-source release is listed for Nebius AI Cloud. Replicate Deployments is open source (Apache-2.0).\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.json, and with the fewest tokens: https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"nebius-ai-cloud\", \"b\": \"replicate-deploy\"}`. From a terminal: `anchor compare nebius-ai-cloud replicate-deploy`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/nebius-ai-cloud.json and https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Other comparisons with Nebius AI Cloud or Replicate Deployments\n\n- [Baseten vs Nebius AI Cloud](https://www.anchorterminal.com/compare/baseten-vs-nebius-ai-cloud.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Beam vs Nebius AI Cloud](https://www.anchorterminal.com/compare/beam-vs-nebius-ai-cloud.md)\n- [Beam vs Replicate Deployments](https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.md)\n- [Cerebrium vs Nebius AI Cloud](https://www.anchorterminal.com/compare/cerebrium-vs-nebius-ai-cloud.md)\n- [Cerebrium vs Replicate Deployments](https://www.anchorterminal.com/compare/cerebrium-vs-replicate-deploy.md)\n- [CoreWeave vs Nebius AI Cloud](https://www.anchorterminal.com/compare/coreweave-vs-nebius-ai-cloud.md)\n- [CoreWeave vs Replicate Deployments](https://www.anchorterminal.com/compare/coreweave-vs-replicate-deploy.md)\n- [Hugging Face Inference Endpoints vs Nebius AI Cloud](https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-nebius-ai-cloud.md)\n- [Hugging Face Inference Endpoints vs Replicate Deployments](https://www.anchorterminal.com/compare/hugging-face-inference-endpoints-vs-replicate-deploy.md)\n- [Hyperbolic vs Nebius AI Cloud](https://www.anchorterminal.com/compare/hyperbolic-vs-nebius-ai-cloud.md)\n- [Hyperbolic vs Replicate Deployments](https://www.anchorterminal.com/compare/hyperbolic-vs-replicate-deploy.md)\n- [Koyeb vs Nebius AI Cloud](https://www.anchorterminal.com/compare/koyeb-vs-nebius-ai-cloud.md)\n- [Koyeb vs Replicate Deployments](https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.md)\n- [Lambda Cloud vs Nebius AI Cloud](https://www.anchorterminal.com/compare/lambda-vs-nebius-ai-cloud.md)\n- [Lambda Cloud vs Replicate Deployments](https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md)\n- [Modal vs Nebius AI Cloud](https://www.anchorterminal.com/compare/modal-vs-nebius-ai-cloud.md)\n- [Modal vs Replicate Deployments](https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.md)\n- [Nebius AI Cloud vs Northflank](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-northflank.md)\n- [Nebius AI Cloud vs Runpod](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-runpod.md)\n- [Nebius AI Cloud vs Thunder Compute](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-thunder-compute.md)\n- [Nebius AI Cloud vs Vast.ai](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-vast-ai.md)\n- [Nebius AI Cloud vs Verda](https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-verda.md)\n- [Northflank vs Replicate Deployments](https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.md)\n- [Replicate Deployments vs Runpod](https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.md)\n- [Replicate Deployments vs Thunder Compute](https://www.anchorterminal.com/compare/replicate-deploy-vs-thunder-compute.md)\n- [Replicate Deployments vs Vast.ai](https://www.anchorterminal.com/compare/replicate-deploy-vs-vast-ai.md)\n- [Replicate Deployments vs Verda](https://www.anchorterminal.com/compare/replicate-deploy-vs-verda.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Nebius AI Cloud vs Replicate Deployments",
        "url": ""
      }
    ],
    "description": "Nebius AI Cloud scores 67.2 (B) on agent readiness against Replicate Deployments's 63.6 (B), and leads in 3 of 7 scored categories. Replicate Deployments leads on reliability, schema \u0026 documentation and payments \u0026 pricing. Both do compute gpu. Category scores, facts, verdicts…",
    "facts": [
      "Nebius AI Cloud B 67.2",
      "Replicate Deployments B 63.6",
      "scores"
    ],
    "h1": "Nebius AI Cloud vs Replicate Deployments",
    "image": "https://www.anchorterminal.com/assets/og/compare-nebius-ai-cloud-vs-replicate-deploy.png",
    "path": "/compare/nebius-ai-cloud-vs-replicate-deploy",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Nebius AI Cloud vs Replicate Deployments for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/nebius-ai-cloud-vs-replicate-deploy"
  },
  "tokens": {
    "markdown": 2750,
    "slim": 730
  },
  "version": 1
}
