{
  "data": {
    "a": {
      "slug": "cartesia-ink-stt",
      "name": "Cartesia Ink",
      "vendor": "Cartesia",
      "vendorUrl": "https://www.cartesia.ai",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Cartesia's hosted speech-to-text API. Ink 2 transcribes live audio in five languages over a WebSocket with built-in turn detection, and the older Ink Whisper model transcribes uploaded files in about 100 languages.",
      "url": "https://www.anchorterminal.com/tools/cartesia-ink-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/cartesia-ink-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/cartesia-ink-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/cartesia-ink-stt.json",
      "repo": "https://github.com/cartesia-ai/cartesia-python",
      "license": "Proprietary hosted service under the Cartesia Terms of Service. The Python and JavaScript SDKs are Apache-2.0",
      "transports": [
        "http",
        "websocket",
        "streamable-http"
      ],
      "remoteUrl": "https://api.cartesia.ai",
      "packages": [
        {
          "registry": "pypi",
          "name": "cartesia"
        },
        {
          "registry": "npm",
          "name": "@cartesia/cartesia-js"
        },
        {
          "registry": "pypi",
          "name": "cartesia-mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve API key (`sk_car_...`) from the Playground at play.cartesia.ai/keys, sent as `Authorization: Bearer` or `X-API-Key`. Browser clients use an access token with an `stt` grant, valid for at most one hour, passed as the `access_token` query parameter on WebSockets. Usage and key-metadata endpoints need a separate admin key (https://docs.cartesia.ai/use-the-api/api-conventions).",
      "pricing": "freemium",
      "pricingNotes": "Billed in credits. `ink-2` costs 3 credits a second of audio on both realtime endpoints, silence included. `ink-whisper` costs 1 credit a second in realtime and 1 credit per 2 seconds in batch (https://docs.cartesia.ai/pricing). Plans are Free ($0, 20,000 credits a month), Pro ($5, 100,000), Startup ($49, 1.25 million) and Scale ($299, 8 million), so an agent's owner can start without a contract. Commercial use starts at Pro (https://www.cartesia.ai/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the docs index, the pricing pages or the STT reference (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 134,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://docs.cartesia.ai/use-the-api/stt/compare-endpoints",
      "llmsTxt": "https://docs.cartesia.ai/llms.txt",
      "openapi": "https://docs.cartesia.ai/latest.yml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "closed-source",
        "streaming",
        "batch",
        "free-tier",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "mcp",
        "status-page",
        "security-txt",
        "soc2",
        "enterprise"
      ],
      "lastRelease": "2026-09-17",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67.9,
        "grade": "B",
        "agentReady": false,
        "rank": 238,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 9,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 74,
          "maintenance": 85,
          "payments": 35,
          "reliability": 73,
          "schema": 89,
          "security": 52,
          "transparency": 67
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "Ink 2 streams transcripts with turn detection built into the model, documented by an AsyncAPI file with typed ranges and structured errors. The terms, last revised 14 June 2024, let Cartesia train on inputs and outputs unless the customer opts out, and zero data retention is an Enterprise setting. Ink 2 has no batch endpoint and no diarisation.",
        "bestFor": "Suited to voice agents that need turn detection and transcription from one model in English, French, Hindi, Japanese or Spanish, and to teams already using Cartesia text-to-speech.",
        "strengths": [
          "`/stt/turns/websocket` emits `turn.start`, `turn.update`, `turn.eager_end`, `turn.resume` and `turn.end`, so no separate voice activity detector is needed",
          "OpenAPI and AsyncAPI files, llms.txt and Markdown twins of every docs page, with enums and numeric ranges on the WebSocket parameters",
          "Structured errors on `Cartesia-Version` 2026-03-01 and later, with `error_code`, `title`, `message`, `request_id` and an optional `doc_url`",
          "Short-lived access tokens (at most one hour) carry an `stt` grant, and admin keys are a separate key type from standard keys",
          "The status page lists Speech to Text in four regions, at 99.97% (US), 99.994% (EU and APAC) and 100% (AU) for July to October 2026"
        ],
        "weaknesses": [
          "The terms let Cartesia train models on inputs and outputs unless otherwise agreed. Opting out is a form in the Playground's data controls",
          "Zero data retention is an Enterprise plan setting, and no retention period for other plans was found in the terms, privacy policy or docs",
          "`ink-2` supports English, French, Hindi, Japanese and Spanish only and has no batch endpoint. `POST /stt` accepts `ink-whisper` only",
          "No diarisation or speaker labels in the reviewed documentation, and realtime input is raw mono PCM with `encoding` and `sample_rate` required",
          "The terms of 14 June 2024 forbid automation software (bots) and any robot or scraper that accesses the Services to collect data, which matters before any probe is run",
          "No SLA was found, and a 429 for exceeding the concurrency limit is documented without a `Retry-After` header or backoff guidance"
        ],
        "agentNotes": [
          "Use `wss://api.cartesia.ai/stt/turns/websocket` with `model=ink-2`, `encoding`, `sample_rate` and `cartesia_version=2026-08-14`. All four are required.",
          "Send raw mono audio in chunks of about 100 ms at the speed it was spoken. Pushing a whole file into the socket can return an internal server error.",
          "Read the final text from `turn.end` only. `transcript` is cumulative within a turn, so joining `turn.update` events duplicates text.",
          "Send `{\"type\": \"close\"}` after the last audio and keep reading until the server closes the socket, or the buffered tail is lost.",
          "Check `encoding` and `sample_rate` against the source before sending. The docs say the server might not return an error when they are wrong."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67.9
          }
        ],
        "editorialScores": {
          "ergonomics": 74,
          "maintenance": 85,
          "payments": 35,
          "reliability": 73,
          "schema": 89,
          "security": 52,
          "transparency": 49
        },
        "provenanceScore": 84
      },
      "connect": {
        "install": "pip install 'cartesia[websockets]'   # or: npm install @cartesia/cartesia-js",
        "claudeCode": "claude mcp add --transport http --scope user cartesia https://mcp.cartesia.ai/mcp"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/cartesia-ink-stt"
      },
      "sameCompany": [
        "cartesia-tts",
        "cartesia-voice-cloning"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Ink 2 realtime, Scale plan",
          "unit": "audio-minute",
          "usd": 0.0067,
          "note": "180 credits a minute at $299 for 8M credits. Cartesia quotes $0.39 an hour"
        },
        {
          "item": "Ink 2 realtime, Startup plan",
          "unit": "audio-minute",
          "usd": 0.0071,
          "note": "180 credits a minute at $49 for 1.25M credits"
        },
        {
          "item": "Ink 2 realtime, Pro plan",
          "unit": "audio-minute",
          "usd": 0.009,
          "note": "180 credits a minute at $5 for 100,000 credits"
        },
        {
          "item": "Ink Whisper realtime, Scale plan",
          "unit": "audio-minute",
          "usd": 0.0022,
          "note": "60 credits a minute at $299 for 8M credits"
        },
        {
          "item": "Ink Whisper batch, Scale plan",
          "unit": "audio-minute",
          "usd": 0.0011,
          "note": "30 credits a minute at $299 for 8M credits"
        }
      ],
      "provenance": {
        "legalEntity": "Cartesia AI, Inc.",
        "domain": "cartesia.ai",
        "domainRegistered": "2023-05-10",
        "endpointOnVendorDomain": true,
        "terms": "https://www.cartesia.ai/legal/terms",
        "privacy": "https://www.cartesia.ai/legal/privacy",
        "statusPage": "https://status.cartesia.ai",
        "changelog": "https://docs.cartesia.ai/changelog/2026",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 84
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/cartesia-ink-stt.json",
      "live": {
        "slug": "cartesia-ink-stt",
        "probe": {
          "target": "https://api.cartesia.ai",
          "method": "get",
          "lastAt": "2026-10-10T03:07:00.809061285Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 126,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 69,
          "p95ms24h": 203,
          "samples24h": 117,
          "samples30d": 117,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cartesia.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T03:01:52.995710279Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "cartesia-ai/cartesia-python",
            "version": "v4.2.0",
            "released": "2026-09-02",
            "seenAt": "2026-10-09T16:44:27.669680758Z"
          },
          {
            "registry": "npm",
            "name": "@cartesia/cartesia-js",
            "version": "4.2.0",
            "seenAt": "2026-10-09T16:44:25.845737197Z"
          },
          {
            "registry": "pypi",
            "name": "cartesia",
            "version": "4.2.0",
            "released": "2026-09-02",
            "seenAt": "2026-10-09T16:44:25.659263821Z"
          },
          {
            "registry": "pypi",
            "name": "cartesia-mcp",
            "version": "0.26.1",
            "released": "2026-10-08",
            "seenAt": "2026-10-09T16:44:26.977578193Z"
          }
        ],
        "githubStars": 134,
        "npmWeekly": 79696,
        "pypiWeekly": 205864,
        "pages": [
          {
            "url": "https://docs.cartesia.ai/changelog/2026",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:36:45.184290806Z",
            "changedAt": "2026-10-07T18:04:49.25388057Z",
            "fingerprint": "405fadad924b"
          },
          {
            "url": "https://docs.cartesia.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:36:47.532132362Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "379a41a8bcf3"
          },
          {
            "url": "https://www.cartesia.ai/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:59.764022121Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "12bd9269bdef"
          },
          {
            "url": "https://www.cartesia.ai/legal/privacy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:55.522982064Z",
            "changedAt": "2026-10-08T18:26:54.204830291Z",
            "fingerprint": "90fbd63881e5"
          },
          {
            "url": "https://www.cartesia.ai/legal/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:48:57.753461207Z",
            "changedAt": "2026-10-08T18:26:56.591248396Z",
            "fingerprint": "2153bed03a23"
          }
        ],
        "updatedAt": "2026-10-10T03:07:00.809061285Z"
      }
    },
    "answer": "Google Cloud Speech-to-Text scores 70.2 (BB) on agent readiness against Cartesia Ink's 67.9 (B), and leads in 3 of 7 scored categories. Cartesia Ink leads on schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.",
    "b": {
      "slug": "google-speech-to-text",
      "name": "Google Cloud Speech-to-Text",
      "vendor": "Google Cloud",
      "vendorUrl": "https://cloud.google.com/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Google Cloud's transcription API.",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "repo": "https://github.com/googleapis/google-cloud-python/tree/main/packages/google-cloud-speech",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://speech.googleapis.com/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-cloud-speech"
        },
        {
          "registry": "npm",
          "name": "@google-cloud/speech"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth 2.0 bearer token from a service account or `gcloud` (Application Default Credentials) on a project with billing and the API turned on. Chirp 3 runs on the `us` and `eu` multi-region endpoints such as `us-speech.googleapis.com`.",
      "pricing": "freemium",
      "pricingNotes": "V2 standard recognition, which covers Chirp 3, is $0.016 a minute to 500,000 minutes a month, then $0.01, $0.008 and $0.004 past 2M. Dynamic batch is $0.003 a minute. Billed per second, per channel. V1 has 60 free minutes a month and charges $0.024 without data logging. Medical models $0.078 (https://cloud.google.com/speech-to-text/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 713013,
        "pypiWeekly": 3703632,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.cloud.google.com/speech-to-text/docs",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "card-required"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.2,
        "grade": "BB",
        "agentReady": true,
        "rank": 159,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 6,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 86
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
        "bestFor": "Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.",
        "strengths": [
          "Audio isn't stored or used for training unless the project opts in to data logging",
          "No Speech-to-Text incident on the Google Cloud status page since 12 June 2025",
          "Dynamic batch at $0.003 a minute, and standard recognition tiers down to $0.004 past 2M minutes",
          "OAuth service accounts with IAM roles, and Cloud Audit Logs",
          "300 concurrent streams per region by default"
        ],
        "weaknesses": [
          "No release note since 2025-11-13",
          "82 of the 111 Chirp 3 locales are preview",
          "Sync requests stop at 1 minute and streams at 5 minutes, and batch reads only from Cloud Storage",
          "No API keys in the documented V2 flow, and the free minutes need a billed project",
          "The quotas page doesn't say what error a breach returns or how to back off"
        ],
        "agentNotes": [
          "Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location",
          "Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings",
          "Downmix stereo unless you need channel labels, since each channel is billed",
          "Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute",
          "Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.2
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 75
        },
        "provenanceScore": 97
      },
      "connect": {
        "install": "pip install google-cloud-speech   # or: npm i @google-cloud/speech",
        "http": "curl -X POST \"https://us-speech.googleapis.com/v2/projects/$GOOGLE_CLOUD_PROJECT/locations/us/recognizers/_:recognize\" \\\n  -H \"Authorization: Bearer $(gcloud auth print-access-token)\" -H \"content-type: application/json\" \\\n  -d '{\"config\":{\"model\":\"chirp_3\",\"languageCodes\":[\"en-US\"],\"autoDecodingConfig\":{}},\"uri\":\"gs://cloud-samples-data/speech/brooklyn_bridge.flac\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/google-speech-to-text"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "gemini-live",
        "google-adk",
        "google-secret-manager",
        "google-cloud-document-ai",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "firebase-cloud-messaging",
        "google-drive-api",
        "gemini-cli",
        "google-search-console",
        "google-ads-api",
        "google-forms",
        "google-sheets-api",
        "gmail-api"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "V2 standard recognition (Chirp 3)",
          "unit": "audio-minute",
          "usd": 0.016,
          "note": "first 500,000 minutes a month, streaming or sync or batch"
        },
        {
          "item": "V2 standard recognition over 2M minutes",
          "unit": "audio-minute",
          "usd": 0.004
        },
        {
          "item": "V2 dynamic batch",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "lower-priority batch"
        },
        {
          "item": "V1 without data logging",
          "unit": "audio-minute",
          "usd": 0.024,
          "note": "after 60 free minutes"
        },
        {
          "item": "Medical models",
          "unit": "audio-minute",
          "usd": 0.078
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "domainNote": "The endpoint is on googleapis.com, Google's API domain. google.com was registered in 1997.",
        "endpointOnVendorDomain": true,
        "terms": "https://cloud.google.com/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://status.cloud.google.com",
        "changelog": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 97
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
      "live": {
        "slug": "google-speech-to-text",
        "probe": {
          "target": "https://speech.googleapis.com/v2",
          "method": "get",
          "lastAt": "2026-10-10T03:07:07.050260684Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 44,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 45,
          "p95ms24h": 88,
          "samples24h": 249,
          "samples30d": 2468,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cloud.google.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-09-30T22:44:37.367865472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "googleapis/google-cloud-python",
            "version": "google-devicesandservices-health-v0.1.4",
            "released": "2026-10-08",
            "seenAt": "2026-10-09T16:56:11.571804166Z"
          },
          {
            "registry": "npm",
            "name": "@google-cloud/speech",
            "version": "8.1.1",
            "seenAt": "2026-10-09T16:56:11.174516519Z"
          },
          {
            "registry": "pypi",
            "name": "google-cloud-speech",
            "version": "2.41.0",
            "released": "2026-10-01",
            "seenAt": "2026-10-09T16:56:10.983522704Z"
          }
        ],
        "githubStars": 5404,
        "npmWeekly": 794903,
        "pypiWeekly": 3248560,
        "securityTxt": {
          "url": "https://google.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2030-04-01T00:00:00z",
          "checkedAt": "2026-10-09T15:39:28.443612679Z"
        },
        "domain": {
          "domain": "google.com",
          "registered": "1997-09-15",
          "source": "https://rdap.verisign.com/com/v1/domain/google.com",
          "checkedAt": "2026-10-04T13:05:50.737985829Z"
        },
        "pages": [
          {
            "url": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:36:59.279441468Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "03f9dac9276b"
          },
          {
            "url": "https://cloud.google.com/speech-to-text/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:33:59.03936004Z",
            "changedAt": "2026-10-08T18:16:21.715117955Z",
            "fingerprint": "8dbbc25a8959"
          },
          {
            "url": "https://cloud.google.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-01T13:11:34.992628421Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6798e0f4fb24"
          }
        ],
        "updatedAt": "2026-10-10T03:07:07.050260684Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Cartesia",
        "b": "Google Cloud",
        "name": "Vendor"
      },
      {
        "a": "https://api.cartesia.ai",
        "b": "https://speech.googleapis.com/v2",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, websocket, Streamable HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "OAuth",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service under the Cartesia Terms of Service. The Python and JavaScript SDKs are Apache-2.0",
        "b": "Apache-2.0 (SDKs)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "no",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-17",
        "b": "2026-09-28",
        "name": "Last release"
      },
      {
        "a": "2024-06-14",
        "b": "2026-09-02",
        "name": "Terms last updated"
      },
      {
        "a": "2024-06-14",
        "b": "2026-10-01",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "yes",
        "name": "Customer content may train models"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "134 stars",
        "b": "713k npm/wk, 3.7M PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Google Cloud Speech-to-Text scores 70.2 (BB) on agent readiness against Cartesia Ink's 67.9 (B), and leads in 3 of 7 scored categories. Cartesia Ink leads on schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.",
        "question": "Which is better for AI agents, Cartesia Ink or Google Cloud Speech-to-Text?"
      },
      {
        "answer": "Cartesia Ink needs an API key. Google Cloud Speech-to-Text uses an OAuth sign-in.",
        "question": "Do Cartesia Ink and Google Cloud Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. Cartesia Ink has a hosted endpoint at https://api.cartesia.ai and Google Cloud Speech-to-Text at https://speech.googleapis.com/v2.",
        "question": "Can an agent call Cartesia Ink and Google Cloud Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Schema \u0026 documentation, 89 against 80",
          "Payments \u0026 pricing, 35 against 20",
          "Maintenance \u0026 community, 85 against 25"
        ],
        "also": null,
        "goodFor": "Suited to voice agents that need turn detection and transcription from one model in English, French, Hindi, Japanese or Spanish, and to teams already using Cartesia text-to-speech.",
        "slug": "cartesia-ink-stt",
        "watchFor": "The terms let Cartesia train models on inputs and outputs unless otherwise agreed. Opting out is a form in the Playground's data controls"
      },
      {
        "aheadOn": [
          "Reliability, 85 against 73",
          "Security \u0026 auth, 95 against 52",
          "Transparency \u0026 trust, 86 against 67"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.",
        "slug": "google-speech-to-text",
        "watchFor": "No release note since 2025-11-13"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-cartesia-ink-stt.json",
        "title": "Amazon Transcribe vs Cartesia Ink",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-cartesia-ink-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.json",
        "title": "Amazon Transcribe vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Cartesia Ink",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-cartesia-ink-stt.json",
        "title": "Azure AI Speech speech-to-text vs Cartesia Ink",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-cartesia-ink-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt.json",
        "title": "Cartesia Ink vs Deepgram Speech-to-Text (Nova-3, Flux)",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-elevenlabs-scribe.json",
        "title": "Cartesia Ink vs ElevenLabs Scribe Speech to Text API",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-elevenlabs-scribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-gladia-stt.json",
        "title": "Cartesia Ink vs Gladia Speech-to-Text API + MCP",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-gladia-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text.json",
        "title": "Cartesia Ink vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe.json",
        "title": "Cartesia Ink vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-rev-ai-stt.json",
        "title": "Cartesia Ink vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt.json",
        "title": "Cartesia Ink vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt.json",
        "title": "Cartesia Ink vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.json",
        "title": "Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.json",
        "title": "Google Cloud Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.json",
        "title": "Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 12,
        "cartesia-ink-stt": 73,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 85,
        "key": "reliability",
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 9,
        "cartesia-ink-stt": 89,
        "edge": "cartesia-ink-stt",
        "google-speech-to-text": 80,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 4,
        "cartesia-ink-stt": 74,
        "edge": "cartesia-ink-stt",
        "google-speech-to-text": 70,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 43,
        "cartesia-ink-stt": 52,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 95,
        "key": "security",
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 15,
        "cartesia-ink-stt": 35,
        "edge": "cartesia-ink-stt",
        "google-speech-to-text": 20,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 60,
        "cartesia-ink-stt": 85,
        "edge": "cartesia-ink-stt",
        "google-speech-to-text": 25,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 19,
        "cartesia-ink-stt": 67,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 86,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Google Cloud Speech-to-Text scores 70.2 (BB) on agent readiness against Cartesia Ink's 67.9 (B), and leads in 3 of 7 scored categories. Cartesia Ink leads on schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community. Both do speech-to-text.",
    "verdicts": {
      "cartesia-ink-stt": "Ink 2 streams transcripts with turn detection built into the model, documented by an AsyncAPI file with typed ranges and structured errors. The terms, last revised 14 June 2024, let Cartesia train on inputs and outputs unless the customer opts out, and zero data retention is an Enterprise setting. Ink 2 has no batch endpoint and no diarisation.",
      "google-speech-to-text": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.min.md"
  },
  "markdown": "Google Cloud Speech-to-Text scores 70.2 (BB) on agent readiness against Cartesia Ink's 67.9 (B), and leads in 3 of 7 scored categories. Cartesia Ink leads on schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community. Both do speech-to-text.\n\n- Cartesia Ink: grade B, 67.9/100, rank #238 of 950. Markdown https://www.anchorterminal.com/tools/cartesia-ink-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/cartesia-ink-stt.json\n- Google Cloud Speech-to-Text: grade BB, 70.2/100, rank #159 of 950. Markdown https://www.anchorterminal.com/tools/google-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### Cartesia Ink (B)\n\nGood for: Suited to voice agents that need turn detection and transcription from one model in English, French, Hindi, Japanese or Spanish, and to teams already using Cartesia text-to-speech.\n\nAhead on:\n- Schema \u0026 documentation, 89 against 80\n- Payments \u0026 pricing, 35 against 20\n- Maintenance \u0026 community, 85 against 25\n\nWatch for: The terms let Cartesia train models on inputs and outputs unless otherwise agreed. Opting out is a form in the Playground's data controls\n\n### Google Cloud Speech-to-Text (BB)\n\nGood for: Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.\n\nAhead on:\n- Reliability, 85 against 73\n- Security \u0026 auth, 95 against 52\n- Transparency \u0026 trust, 86 against 67\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: No release note since 2025-11-13\n\n\n## Score by category\n\n| Category | Weight | Cartesia Ink | Google Cloud Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 73 | 85 | Google Cloud Speech-to-Text +12 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 89 | 80 | Cartesia Ink +9 |\n| Agent ergonomics | 13% (16.2 this run) | 74 | 70 | Cartesia Ink +4 |\n| Security \u0026 auth | 14% (17.5 this run) | 52 | 95 | Google Cloud Speech-to-Text +43 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 35 | 20 | Cartesia Ink +15 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 85 | 25 | Cartesia Ink +60 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 67 | 86 | Google Cloud Speech-to-Text +19 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **67.9 · B** | **70.2 · BB** | |\n\n## Facts side by side\n\n| Fact | Cartesia Ink | Google Cloud Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Cartesia | Google Cloud |\n| Hosted endpoint | `https://api.cartesia.ai` | `https://speech.googleapis.com/v2` |\n| Transports | HTTP, websocket, Streamable HTTP | HTTP |\n| Auth | API key | OAuth |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | Proprietary hosted service under the Cartesia Terms of Service. The Python and JavaScript SDKs are Apache-2.0 | Apache-2.0 (SDKs) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| Last release | 2026-09-17 | 2026-09-28 |\n| Terms last updated | 2024-06-14 | 2026-09-02 |\n| Privacy policy last updated | 2024-06-14 | 2026-10-01 |\n| Customer content may train models | yes, with an opt-out | yes |\n| Terms restrict automated access | yes | not found in the text |\n| Terms restrict benchmarking | yes | not found in the text |\n| Terms or service can change without notice | not found in the text | not found in the text |\n| Arbitration or class-action waiver | yes | not found in the text |\n| Popularity | 134 stars | 713k npm/wk, 3.7M PyPI/wk |\n| Agent reviews | none | 3/5 (2) |\n\n## Verdicts\n\n**Cartesia Ink.** Ink 2 streams transcripts with turn detection built into the model, documented by an AsyncAPI file with typed ranges and structured errors. The terms, last revised 14 June 2024, let Cartesia train on inputs and outputs unless the customer opts out, and zero data retention is an Enterprise setting. Ink 2 has no batch endpoint and no diarisation.\n\n**Google Cloud Speech-to-Text.** Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.\n\n## Before you call either\n\n### Cartesia Ink\n\n1. Use `wss://api.cartesia.ai/stt/turns/websocket` with `model=ink-2`, `encoding`, `sample_rate` and `cartesia_version=2026-08-14`. All four are required.\n2. Send raw mono audio in chunks of about 100 ms at the speed it was spoken. Pushing a whole file into the socket can return an internal server error.\n3. Read the final text from `turn.end` only. `transcript` is cumulative within a turn, so joining `turn.update` events duplicates text.\n4. Send `{\"type\": \"close\"}` after the last audio and keep reading until the server closes the socket, or the buffered tail is lost.\n5. Check `encoding` and `sample_rate` against the source before sending. The docs say the server might not return an error when they are wrong.\n\n### Google Cloud Speech-to-Text\n\n1. Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location\n2. Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings\n3. Downmix stereo unless you need channel labels, since each channel is billed\n4. Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute\n5. Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval\n\n## Questions\n\n### Which is better for AI agents, Cartesia Ink or Google Cloud Speech-to-Text?\n\nGoogle Cloud Speech-to-Text scores 70.2 (BB) on agent readiness against Cartesia Ink's 67.9 (B), and leads in 3 of 7 scored categories. Cartesia Ink leads on schema \u0026 documentation, payments \u0026 pricing and maintenance \u0026 community.\n\n### Do Cartesia Ink and Google Cloud Speech-to-Text need an API key?\n\nCartesia Ink needs an API key. Google Cloud Speech-to-Text uses an OAuth sign-in.\n\n### Can an agent call Cartesia Ink and Google Cloud Speech-to-Text without installing anything?\n\nYes. Cartesia Ink has a hosted endpoint at https://api.cartesia.ai and Google Cloud Speech-to-Text at https://speech.googleapis.com/v2.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"cartesia-ink-stt\", \"b\": \"google-speech-to-text\"}`. From a terminal: `anchor compare cartesia-ink-stt google-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/cartesia-ink-stt.json and https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n\n## Other comparisons with Cartesia Ink or Google Cloud Speech-to-Text\n\n- [Amazon Transcribe vs Cartesia Ink](https://www.anchorterminal.com/compare/amazon-transcribe-vs-cartesia-ink-stt.md)\n- [Amazon Transcribe vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Cartesia Ink](https://www.anchorterminal.com/compare/assemblyai-stt-vs-cartesia-ink-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Cartesia Ink](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-cartesia-ink-stt.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Cartesia Ink vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-deepgram-stt.md)\n- [Cartesia Ink vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-elevenlabs-scribe.md)\n- [Cartesia Ink vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-gladia-stt.md)\n- [Cartesia Ink vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text.md)\n- [Cartesia Ink vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-mistral-voxtral-transcribe.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-rev-ai-stt.md)\n- [Cartesia Ink vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-soniox-stt.md)\n- [Cartesia Ink vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-speechmatics-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Cartesia Ink vs Google Cloud Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Google Cloud Speech-to-Text scores 70.2 (BB) to Cartesia Ink's 67.9 (B) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Cartesia Ink B 67.9",
      "Google Cloud Speech-to-Text BB 70.2",
      "scores"
    ],
    "h1": "Cartesia Ink vs Google Cloud Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-cartesia-ink-stt-vs-google-speech-to-text.png",
    "path": "/compare/cartesia-ink-stt-vs-google-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Cartesia Ink vs Google Cloud Speech-to-Text for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text"
  },
  "tokens": {
    "markdown": 2750,
    "slim": 730
  },
  "version": 1
}
