{
  "data": {
    "a": {
      "slug": "lambda",
      "name": "Lambda Cloud",
      "vendor": "Lambda",
      "vendorUrl": "https://lambda.ai",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "On-demand GPU virtual machines and clusters, with an API for provisioning compute and persistent storage.",
      "url": "https://www.anchorterminal.com/tools/lambda",
      "markdownUrl": "https://www.anchorterminal.com/tools/lambda.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/lambda.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/lambda.json",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://cloud.lambda.ai/api/v1",
      "packages": [],
      "auth": "api-key",
      "authNotes": "API key created at cloud.lambda.ai/api-keys and sent as `Authorization: Bearer`. HTTP Basic with the key as the username (`curl -u 'KEY:'`) still works as a legacy option. SSH keys registered in the account are injected into launched instances.",
      "pricing": "usage",
      "pricingNotes": "On-demand, per GPU an hour, Tesla V100 16 GB $0.79, A100 40 GB $1.99, A100 80 GB $2.79, H100 SXM 80 GB $3.99, B200 180 GB $6.69. 1-Click Clusters of HGX B200 are quoted per GPU-hour at $9.86 for 16 GPUs, $9.36 for 64 and $8.87 for 256 or more on two-week to one-year commitments. Instances bill in one-minute increments from the moment they pass health checks until you terminate them, invoiced weekly. Filesystems bill per GB used a month in one-hour increments. No free tier (https://lambda.ai/pricing, https://docs.lambda.ai/public-cloud/billing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.lambda.ai/public-cloud/on-demand/",
      "openapi": "https://docs.lambda.ai/api/cloud/spec.json",
      "capabilities": [
        "compute.gpu",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "openapi",
        "enterprise"
      ],
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 50.1,
        "grade": "D",
        "agentReady": false,
        "rank": 363,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 7,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 59
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.",
        "strengths": [
          "H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing",
          "OpenAPI 3.1 spec with documented error codes, a `suggestion` field and cursor pagination",
          "Audit events endpoint filterable by time and resource type",
          "Published rate limits, one request a second and one launch every 12 seconds",
          "Trust portal with SOC 2 Type 2 and four ISO certifications, and a named subprocessor list"
        ],
        "weaknesses": [
          "No scale to zero, autoscaling or endpoints; an idle VM bills until terminated",
          "Ten status incidents in two months, including a two-day regional outage in August 2026",
          "No changelog, no llms.txt and no official SDK",
          "API keys have no scopes and launch has no idempotency key",
          "security.txt expired on 1 June 2026 and no free tier"
        ],
        "agentNotes": [
          "Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch",
          "Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`",
          "Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change",
          "List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine",
          "Terminate the instance in a `finally` block; billing runs by the minute until you do"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 50.1
          }
        ],
        "editorialScores": {
          "ergonomics": 63,
          "maintenance": 5,
          "payments": 20,
          "reliability": 50,
          "schema": 69,
          "security": 60,
          "transparency": 48
        },
        "provenanceScore": 70
      },
      "connect": {
        "http": "curl \"https://cloud.lambda.ai/api/v1/instance-types\" -H \"Authorization: Bearer $LAMBDA_API_KEY\""
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/lambda"
      },
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 3.99,
          "note": "Billed per minute"
        },
        {
          "item": "B200 180 GB",
          "unit": "gpu-hour",
          "usd": 6.69
        },
        {
          "item": "A100 SXM 80 GB",
          "unit": "gpu-hour",
          "usd": 2.79
        },
        {
          "item": "A100 SXM 40 GB",
          "unit": "gpu-hour",
          "usd": 1.99
        },
        {
          "item": "Tesla V100 16 GB",
          "unit": "gpu-hour",
          "usd": 0.79
        },
        {
          "item": "HGX B200 1-Click Cluster, 16 GPUs",
          "unit": "gpu-hour",
          "usd": 9.86,
          "note": "Two-week to one-year commitment; $8.87 at 256 GPUs or more"
        }
      ],
      "provenance": {
        "legalEntity": "Lambda, Inc.",
        "domain": "lambda.ai",
        "domainRegistered": "",
        "domainNote": "Lambda moved from lambdalabs.com, registered 2008-05-29, to lambda.ai. The .ai registry's RDAP server rate-limited our lookup of the new domain.",
        "endpointOnVendorDomain": true,
        "terms": "https://lambda.ai/legal/terms-of-service",
        "privacy": "https://lambda.ai/legal/privacy-policy",
        "statusPage": "https://status.lambda.ai",
        "changelog": "",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "notes": [
          "Terms dated August 2025 and the privacy policy of 1 January 2026 name Lambda, Inc., 2510 Zanker Road, San Jose, California.",
          "The API runs on cloud.lambda.ai, a subdomain of the vendor domain.",
          "security.txt expired 2026-06-01 and has no Policy field.",
          "No public changelog found for the cloud. The OpenAPI spec reports version 1.10.0."
        ],
        "score": 70
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/lambda.json",
      "live": {
        "slug": "lambda",
        "probe": {
          "target": "https://cloud.lambda.ai/api/v1",
          "method": "get",
          "lastAt": "2026-10-04T21:48:30.427410181Z",
          "lastOk": true,
          "lastStatus": 403,
          "lastMs": 574,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 561,
          "p95ms24h": 901,
          "samples24h": 272,
          "samples30d": 875,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 247,
              "ok": 247
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.lambda.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T21:40:11.367004648Z"
        },
        "securityTxt": {
          "url": "https://lambda.ai/.well-known/security.txt",
          "state": "expired",
          "expires": "2026-06-01T16:00:00Z",
          "checkedAt": "2026-10-04T15:15:58.764976134Z"
        },
        "domain": {
          "domain": "lambda.ai",
          "registered": "2017-12-16",
          "source": "https://rdap.identitydigital.services/rdap/domain/lambda.ai",
          "checkedAt": "2026-10-04T13:04:37.89683091Z"
        },
        "pages": [
          {
            "url": "https://lambda.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:18.84160567Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ad7e156559d1"
          },
          {
            "url": "https://lambda.ai/legal/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:14.766678886Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "921eec75dd83"
          },
          {
            "url": "https://lambda.ai/legal/terms-of-service",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:45:16.817640272Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "40e40aee2e7c"
          }
        ],
        "updatedAt": "2026-10-04T21:48:30.427410181Z"
      }
    },
    "b": {
      "slug": "replicate-deploy",
      "name": "Replicate Deployments",
      "vendor": "Replicate",
      "vendorUrl": "https://replicate.com",
      "kind": "http-api",
      "category": "gpu-compute",
      "summary": "Replicate's service for deploying and running custom models.",
      "url": "https://www.anchorterminal.com/tools/replicate-deploy",
      "markdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/replicate-deploy.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json",
      "repo": "https://github.com/replicate/cog",
      "license": "Apache-2.0",
      "transports": [
        "http",
        "sse",
        "stdio"
      ],
      "remoteUrl": "https://api.replicate.com/v1",
      "packages": [
        {
          "registry": "npm",
          "name": "replicate"
        },
        {
          "registry": "pypi",
          "name": "replicate"
        },
        {
          "registry": "npm",
          "name": "replicate-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API token on every call to api.replicate.com. `cog push` uses the same token to upload a model image. The hosted MCP at https://mcp.replicate.com/sse asks for the token in a browser flow and holds it for the client; the local `replicate-mcp` package reads `REPLICATE_API_TOKEN`.",
      "pricing": "usage",
      "pricingNotes": "Private models and deployments bill per second for the whole time an instance is up, set-up and idle included, from prepaid credit or monthly in arrears. CPU $0.000100 a second ($0.36 an hour), T4 $0.000225 ($0.81), L40S $0.000975 ($3.51), A100 80 GB $0.001400 ($5.04), H100 $0.001525 ($5.49), 2x L40S $0.001950 ($7.02), 2x A100 $0.002800 ($10.08). 2x H100 ($10.98), 4x and 8x L40S, A100 and H100 up to $43.92 an hour need a committed-spend contract. Fast-booting fine-tunes bill only while active. Public models bill only active time and not failures (https://replicate.com/pricing, https://replicate.com/docs/topics/billing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 9500,
        "npmWeekly": 634116,
        "pypiWeekly": 386704,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://replicate.com/docs/topics/deployments",
      "llmsTxt": "https://replicate.com/docs/llms.txt",
      "openapi": "https://api.replicate.com/openapi.json",
      "capabilities": [
        "compute.gpu",
        "compute.endpoints",
        "compute.serverless",
        "compute.containers"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "mcp",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "async-jobs",
        "webhooks",
        "open-source"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 63.7,
        "grade": "B",
        "agentReady": false,
        "rank": 197,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 3,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 80
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.",
        "strengths": [
          "OpenAPI file, llms.txt and an MCP server with a two-tool code mode",
          "Deployment min and max instances settable over the API, 0 allowed",
          "API prediction data deleted after one hour by default",
          "Published limits, 600 prediction creates and 3,000 other calls a minute",
          "Leaked tokens found on GitHub are disabled automatically"
        ],
        "weaknesses": [
          "Private instances bill set-up and idle time, H100 at $5.49 an hour",
          "API tokens have no scopes, expiry or audit log",
          "Changelog silent since 21 April 2026",
          "Only T4, L40S, A100 and H100, and more than 2 GPUs needs a committed-spend contract",
          "Two September 2026 incidents ran 15 and 20 hours, both marked minor"
        ],
        "agentNotes": [
          "List `GET /v1/hardware` first and use the returned `sku` in the deployment body",
          "Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not",
          "Send `Prefer: wait` on deployment predictions to block instead of polling",
          "Copy outputs within an hour; API prediction data is deleted after that",
          "Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 63.7
          }
        ],
        "editorialScores": {
          "ergonomics": 68,
          "maintenance": 70,
          "payments": 30,
          "reliability": 75,
          "schema": 85,
          "security": 40,
          "transparency": 69
        },
        "provenanceScore": 90
      },
      "connect": {
        "install": "pip install cog replicate",
        "http": "curl -X POST \"https://api.replicate.com/v1/deployments/$REPLICATE_OWNER/my-deployment/predictions\" \\\n  -H \"Authorization: Bearer $REPLICATE_API_TOKEN\" -H \"Content-Type: application/json\" -H \"Prefer: wait\" \\\n  -d '{\"input\":{\"prompt\":\"hello\"}}'",
        "claudeCode": "claude mcp add replicate https://mcp.replicate.com/sse --transport sse --scope user",
        "config": {
          "mcpServers": {
            "replicate": {
              "args": [
                "-y",
                "replicate-mcp"
              ],
              "command": "npx",
              "env": {
                "REPLICATE_API_TOKEN": "${REPLICATE_API_TOKEN}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/compute.gpu",
        "tool": "https://letme.dev/replicate-deploy"
      },
      "sameCompany": [
        "replicate-image",
        "replicate-musicgen"
      ],
      "area": "models",
      "unitPrices": [
        {
          "item": "H100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.49,
          "note": "$0.001525 a second, including set-up and idle"
        },
        {
          "item": "A100 80 GB",
          "unit": "gpu-hour",
          "usd": 5.04,
          "note": "$0.001400 a second"
        },
        {
          "item": "L40S 48 GB",
          "unit": "gpu-hour",
          "usd": 3.51,
          "note": "$0.000975 a second"
        },
        {
          "item": "T4 16 GB",
          "unit": "gpu-hour",
          "usd": 0.81,
          "note": "$0.000225 a second"
        }
      ],
      "provenance": {
        "legalEntity": "Replicate, LLC",
        "domain": "replicate.com",
        "domainRegistered": "1998-05-26",
        "domainNote": "replicate.com was registered in 1998, long before Replicate the company existed.",
        "endpointOnVendorDomain": true,
        "terms": "https://replicate.com/terms",
        "privacy": "https://replicate.com/privacy",
        "statusPage": "https://replicatestatus.com",
        "changelog": "https://replicate.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Terms last updated 2026-04-01 name Replicate, LLC as the contracting party.",
          "replicatestatus.com redirects to Cloudflare's status page filtered to Replicate.",
          "Replicate's hosted image and music models are listed separately under image generation and music generation."
        ],
        "score": 90
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/replicate-deploy.json",
      "live": {
        "slug": "replicate-deploy",
        "probe": {
          "target": "https://api.replicate.com/v1",
          "method": "get",
          "lastAt": "2026-10-04T21:48:35.132540284Z",
          "lastOk": true,
          "lastStatus": 401,
          "lastMs": 143,
          "lastNote": "asks for credentials",
          "authRequired": true,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 139,
          "p95ms24h": 327,
          "samples24h": 272,
          "samples30d": 875,
          "days": [
            {
              "date": "2026-10-01",
              "probes": 109,
              "ok": 109
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 247,
              "ok": 247
            }
          ]
        },
        "vendorStatus": {
          "page": "https://replicatestatus.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-04T21:40:25.933184888Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "replicate/cog",
            "version": "v0.23.0",
            "released": "2026-09-22",
            "seenAt": "2026-10-04T16:38:03.363821386Z"
          },
          {
            "registry": "npm",
            "name": "replicate",
            "version": "1.4.0",
            "seenAt": "2026-10-04T16:38:01.019271447Z"
          },
          {
            "registry": "npm",
            "name": "replicate-mcp",
            "version": "0.9.0",
            "seenAt": "2026-10-04T16:38:03.126836564Z"
          },
          {
            "registry": "pypi",
            "name": "replicate",
            "version": "1.0.7",
            "released": "2025-05-27",
            "seenAt": "2026-10-04T16:38:01.981507625Z"
          }
        ],
        "githubStars": 9486,
        "npmWeekly": 705916,
        "pypiWeekly": 374867,
        "securityTxt": {
          "url": "https://replicate.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:45.229571506Z"
        },
        "llmsTxt": {
          "url": "https://replicate.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:18:09.902788266Z"
        },
        "domain": {
          "domain": "replicate.com",
          "registered": "1998-05-26",
          "source": "https://rdap.verisign.com/com/v1/domain/replicate.com",
          "checkedAt": "2026-10-04T13:07:04.742407865Z"
        },
        "pages": [
          {
            "url": "https://replicate.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-04T15:47:17.68615995Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "490f4836aca3"
          },
          {
            "url": "https://replicate.com/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:19.974502738Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3f1305f154be"
          },
          {
            "url": "https://replicate.com/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:22.274270654Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "8e299fbc64eb"
          },
          {
            "url": "https://replicate.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:47:23.877388602Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ea48efe3382b"
          }
        ],
        "updatedAt": "2026-10-04T21:48:35.132540284Z"
      }
    },
    "summary": "Replicate Deployments has a score of 63.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 65 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy",
    "json": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.md",
    "slim": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy.min.md"
  },
  "markdown": "Replicate Deployments has a score of 63.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 65 points.\n\n- Lambda Cloud: grade D, 50.1/100, rank #363 of 452. Markdown https://www.anchorterminal.com/tools/lambda.md · JSON https://www.anchorterminal.com/api/v1/tools/lambda.json\n- Replicate Deployments: grade B, 63.7/100, rank #197 of 452. Markdown https://www.anchorterminal.com/tools/replicate-deploy.md · JSON https://www.anchorterminal.com/api/v1/tools/replicate-deploy.json\n\n## Which one, for what\n\nPick Lambda Cloud for security \u0026 auth (+20).\n\nPick Replicate Deployments for reliability (+25), schema \u0026 documentation (+16), agent ergonomics (+5), payments \u0026 pricing (+10), maintenance \u0026 community (+65), transparency \u0026 trust (+21).\n\n## Score by category\n\n| Category | Weight | Lambda Cloud | Replicate Deployments | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 50 | 75 | Replicate Deployments +25 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 69 | 85 | Replicate Deployments +16 |\n| Agent ergonomics | 13% (16.2 this run) | 63 | 68 | Replicate Deployments +5 |\n| Security \u0026 auth | 14% (17.5 this run) | 60 | 40 | Lambda Cloud +20 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 30 | Replicate Deployments +10 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 5 | 70 | Replicate Deployments +65 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 59 | 80 | Replicate Deployments +21 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **50.1 · D** | **63.7 · B** | |\n\n## Facts side by side\n\n| Fact | Lambda Cloud | Replicate Deployments |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Lambda | Replicate |\n| Hosted endpoint | `https://cloud.lambda.ai/api/v1` | `https://api.replicate.com/v1` |\n| Transports | HTTP | HTTP, SSE (legacy), stdio |\n| Auth | API key | API key |\n| Pricing | Pay per use | Pay per use |\n| x402 | no | no |\n| Licence | none | Apache-2.0 |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| MCP registry | not listed | not listed |\n| Last release | none | 2026-09-22 |\n| Popularity | none | 9.5k stars, 634k npm/wk, 387k PyPI/wk |\n| Agent reviews | 3/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**Lambda Cloud.** H100 SXM at $3.99 and B200 at $6.69 an hour with per-minute billing. No scale to zero, autoscaling or endpoints; an idle VM bills until terminated.\n\n**Replicate Deployments.** OpenAPI file, llms.txt and an MCP server with a two-tool code mode. Private instances bill set-up and idle time, H100 at $5.49 an hour.\n\n## Before you call either\n\n### Lambda Cloud\n\n1. Call `GET /instance-types` first and read `regions_with_capacity_available` before trying to launch\n2. Space launch calls 12 seconds apart; a sixth in a minute returns 429 with `global/rate-limited`\n3. Branch on the error `code`, not the `message` or `suggestion`, which Lambda says may change\n4. List instances before retrying a failed launch, since there's no idempotency key and a retry can start a second machine\n5. Terminate the instance in a `finally` block; billing runs by the minute until you do\n\n### Replicate Deployments\n\n1. List `GET /v1/hardware` first and use the returned `sku` in the deployment body\n2. Set `min_instances` to 0 for bursty work; a warm H100 bills $5.49 an hour whether called or not\n3. Send `Prefer: wait` on deployment predictions to block instead of polling\n4. Copy outputs within an hour; API prediction data is deleted after that\n5. Wait for the reset time in the 429 body before retrying; prediction creates cap at 600 a minute\n\n## Other comparisons with Lambda Cloud or Replicate Deployments\n\n- [Baseten vs Lambda Cloud](https://www.anchorterminal.com/compare/baseten-vs-lambda.md)\n- [Baseten vs Replicate Deployments](https://www.anchorterminal.com/compare/baseten-vs-replicate-deploy.md)\n- [Beam vs Lambda Cloud](https://www.anchorterminal.com/compare/beam-vs-lambda.md)\n- [Beam vs Replicate Deployments](https://www.anchorterminal.com/compare/beam-vs-replicate-deploy.md)\n- [Koyeb vs Lambda Cloud](https://www.anchorterminal.com/compare/koyeb-vs-lambda.md)\n- [Koyeb vs Replicate Deployments](https://www.anchorterminal.com/compare/koyeb-vs-replicate-deploy.md)\n- [Lambda Cloud vs Modal](https://www.anchorterminal.com/compare/lambda-vs-modal.md)\n- [Lambda Cloud vs Northflank](https://www.anchorterminal.com/compare/lambda-vs-northflank.md)\n- [Lambda Cloud vs Runpod](https://www.anchorterminal.com/compare/lambda-vs-runpod.md)\n- [Modal vs Replicate Deployments](https://www.anchorterminal.com/compare/modal-vs-replicate-deploy.md)\n- [Northflank vs Replicate Deployments](https://www.anchorterminal.com/compare/northflank-vs-replicate-deploy.md)\n- [Replicate Deployments vs Runpod](https://www.anchorterminal.com/compare/replicate-deploy-vs-runpod.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Lambda Cloud vs Replicate Deployments",
        "url": ""
      }
    ],
    "description": "Replicate Deployments has a score of 63.7 (B) against Lambda Cloud's 50.1 (D). Both do compute gpu. The largest gap is maintenance \u0026 community, 65 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Lambda Cloud D 50.1",
      "Replicate Deployments B 63.7",
      "scores"
    ],
    "h1": "Lambda Cloud vs Replicate Deployments",
    "image": "https://www.anchorterminal.com/assets/og/compare-lambda-vs-replicate-deploy.png",
    "path": "/compare/lambda-vs-replicate-deploy",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Lambda Cloud vs Replicate Deployments for AI agents, D 50.1 vs B 63.7",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/lambda-vs-replicate-deploy"
  },
  "tokens": {
    "markdown": 1450,
    "slim": 380
  },
  "version": 1
}
