{
  "data": {
    "a": {
      "slug": "google-speech-to-text",
      "name": "Google Cloud Speech-to-Text",
      "vendor": "Google Cloud",
      "vendorUrl": "https://cloud.google.com/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Google Cloud's transcription API.",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "repo": "https://github.com/googleapis/google-cloud-python/tree/main/packages/google-cloud-speech",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://speech.googleapis.com/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-cloud-speech"
        },
        {
          "registry": "npm",
          "name": "@google-cloud/speech"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth 2.0 bearer token from a service account or `gcloud` (Application Default Credentials) on a project with billing and the API turned on. Chirp 3 runs on the `us` and `eu` multi-region endpoints such as `us-speech.googleapis.com`.",
      "pricing": "freemium",
      "pricingNotes": "V2 standard recognition, which covers Chirp 3, is $0.016 a minute to 500,000 minutes a month, then $0.01, $0.008 and $0.004 past 2M. Dynamic batch is $0.003 a minute. Billed per second, per channel. V1 has 60 free minutes a month and charges $0.024 without data logging. Medical models $0.078 (https://cloud.google.com/speech-to-text/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 713013,
        "pypiWeekly": 3703632,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.cloud.google.com/speech-to-text/docs",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "card-required"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.2,
        "grade": "BB",
        "agentReady": true,
        "rank": 159,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 6,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 86
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
        "bestFor": "Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.",
        "strengths": [
          "Audio isn't stored or used for training unless the project opts in to data logging",
          "No Speech-to-Text incident on the Google Cloud status page since 12 June 2025",
          "Dynamic batch at $0.003 a minute, and standard recognition tiers down to $0.004 past 2M minutes",
          "OAuth service accounts with IAM roles, and Cloud Audit Logs",
          "300 concurrent streams per region by default"
        ],
        "weaknesses": [
          "No release note since 2025-11-13",
          "82 of the 111 Chirp 3 locales are preview",
          "Sync requests stop at 1 minute and streams at 5 minutes, and batch reads only from Cloud Storage",
          "No API keys in the documented V2 flow, and the free minutes need a billed project",
          "The quotas page doesn't say what error a breach returns or how to back off"
        ],
        "agentNotes": [
          "Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location",
          "Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings",
          "Downmix stereo unless you need channel labels, since each channel is billed",
          "Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute",
          "Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.2
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 75
        },
        "provenanceScore": 97
      },
      "connect": {
        "install": "pip install google-cloud-speech   # or: npm i @google-cloud/speech",
        "http": "curl -X POST \"https://us-speech.googleapis.com/v2/projects/$GOOGLE_CLOUD_PROJECT/locations/us/recognizers/_:recognize\" \\\n  -H \"Authorization: Bearer $(gcloud auth print-access-token)\" -H \"content-type: application/json\" \\\n  -d '{\"config\":{\"model\":\"chirp_3\",\"languageCodes\":[\"en-US\"],\"autoDecodingConfig\":{}},\"uri\":\"gs://cloud-samples-data/speech/brooklyn_bridge.flac\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/google-speech-to-text"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "gemini-live",
        "google-adk",
        "google-secret-manager",
        "google-cloud-document-ai",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "firebase-cloud-messaging",
        "google-drive-api",
        "gemini-cli",
        "google-search-console",
        "google-ads-api",
        "google-forms",
        "google-sheets-api",
        "gmail-api"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "V2 standard recognition (Chirp 3)",
          "unit": "audio-minute",
          "usd": 0.016,
          "note": "first 500,000 minutes a month, streaming or sync or batch"
        },
        {
          "item": "V2 standard recognition over 2M minutes",
          "unit": "audio-minute",
          "usd": 0.004
        },
        {
          "item": "V2 dynamic batch",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "lower-priority batch"
        },
        {
          "item": "V1 without data logging",
          "unit": "audio-minute",
          "usd": 0.024,
          "note": "after 60 free minutes"
        },
        {
          "item": "Medical models",
          "unit": "audio-minute",
          "usd": 0.078
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "domainNote": "The endpoint is on googleapis.com, Google's API domain. google.com was registered in 1997.",
        "endpointOnVendorDomain": true,
        "terms": "https://cloud.google.com/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://status.cloud.google.com",
        "changelog": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 97
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
      "live": {
        "slug": "google-speech-to-text",
        "probe": {
          "target": "https://speech.googleapis.com/v2",
          "method": "get",
          "lastAt": "2026-10-10T02:54:56.717452741Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 37,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 45,
          "p95ms24h": 88,
          "samples24h": 249,
          "samples30d": 2466,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 250,
              "ok": 250
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cloud.google.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-09-30T22:44:37.367865472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "googleapis/google-cloud-python",
            "version": "google-devicesandservices-health-v0.1.4",
            "released": "2026-10-08",
            "seenAt": "2026-10-09T16:56:11.571804166Z"
          },
          {
            "registry": "npm",
            "name": "@google-cloud/speech",
            "version": "8.1.1",
            "seenAt": "2026-10-09T16:56:11.174516519Z"
          },
          {
            "registry": "pypi",
            "name": "google-cloud-speech",
            "version": "2.41.0",
            "released": "2026-10-01",
            "seenAt": "2026-10-09T16:56:10.983522704Z"
          }
        ],
        "githubStars": 5404,
        "npmWeekly": 794903,
        "pypiWeekly": 3248560,
        "securityTxt": {
          "url": "https://google.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2030-04-01T00:00:00z",
          "checkedAt": "2026-10-09T15:39:28.443612679Z"
        },
        "domain": {
          "domain": "google.com",
          "registered": "1997-09-15",
          "source": "https://rdap.verisign.com/com/v1/domain/google.com",
          "checkedAt": "2026-10-04T13:05:50.737985829Z"
        },
        "pages": [
          {
            "url": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:36:59.279441468Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "03f9dac9276b"
          },
          {
            "url": "https://cloud.google.com/speech-to-text/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-09T18:33:59.03936004Z",
            "changedAt": "2026-10-08T18:16:21.715117955Z",
            "fingerprint": "8dbbc25a8959"
          },
          {
            "url": "https://cloud.google.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-01T13:11:34.992628421Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6798e0f4fb24"
          }
        ],
        "updatedAt": "2026-10-10T02:54:56.717452741Z"
      }
    },
    "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Google Cloud Speech-to-Text's 70.2 (BB), and leads in 3 of 7 scored categories. Google Cloud Speech-to-Text leads on reliability, security \u0026 auth and transparency \u0026 trust.",
    "b": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T02:55:04.59384396Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 120,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 138,
          "p95ms24h": 172,
          "samples24h": 115,
          "samples30d": 115,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 30,
              "ok": 30
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T02:50:38.239143757Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T02:55:04.59384396Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Google Cloud",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://speech.googleapis.com/v2",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "OAuth",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Apache-2.0 (SDKs)",
        "b": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "no",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-28",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "2026-09-02",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2026-10-01",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "713k npm/wk, 3.7M PyPI/wk",
        "b": "32k stars",
        "name": "Popularity"
      },
      {
        "a": "3/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Google Cloud Speech-to-Text's 70.2 (BB), and leads in 3 of 7 scored categories. Google Cloud Speech-to-Text leads on reliability, security \u0026 auth and transparency \u0026 trust.",
        "question": "Which is better for AI agents, Google Cloud Speech-to-Text or OpenAI Speech to Text?"
      },
      {
        "answer": "Google Cloud Speech-to-Text uses an OAuth sign-in. OpenAI Speech to Text needs an API key.",
        "question": "Do Google Cloud Speech-to-Text and OpenAI Speech to Text need an API key?"
      },
      {
        "answer": "Yes. Google Cloud Speech-to-Text has a hosted endpoint at https://speech.googleapis.com/v2 and OpenAI Speech to Text at https://api.openai.com/v1.",
        "question": "Can an agent call Google Cloud Speech-to-Text and OpenAI Speech to Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 85 against 80",
          "Security \u0026 auth, 95 against 86",
          "Transparency \u0026 trust, 86 against 61"
        ],
        "also": null,
        "goodFor": "Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.",
        "slug": "google-speech-to-text",
        "watchFor": "No release note since 2025-11-13"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 88 against 80",
          "Agent ergonomics, 78 against 70",
          "Maintenance \u0026 community, 75 against 25"
        ],
        "also": null,
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.json",
        "title": "Amazon Transcribe vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.json",
        "title": "Cartesia Ink vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.json",
        "title": "Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.json",
        "title": "Google Cloud Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.json",
        "title": "Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 5,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 85,
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 8,
        "edge": "openai-speech-to-text",
        "google-speech-to-text": 80,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "weight": 13
      },
      {
        "by": 8,
        "edge": "openai-speech-to-text",
        "google-speech-to-text": 70,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "weight": 13
      },
      {
        "by": 9,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 95,
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "weight": 14
      },
      {
        "by": 0,
        "edge": "",
        "google-speech-to-text": 20,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 50,
        "edge": "openai-speech-to-text",
        "google-speech-to-text": 25,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "weight": 7
      },
      {
        "by": 25,
        "edge": "google-speech-to-text",
        "google-speech-to-text": 86,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Google Cloud Speech-to-Text's 70.2 (BB), and leads in 3 of 7 scored categories. Google Cloud Speech-to-Text leads on reliability, security \u0026 auth and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "google-speech-to-text": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.min.md"
  },
  "markdown": "OpenAI Speech to Text scores 72.4 (BB) on agent readiness against Google Cloud Speech-to-Text's 70.2 (BB), and leads in 3 of 7 scored categories. Google Cloud Speech-to-Text leads on reliability, security \u0026 auth and transparency \u0026 trust. Both do speech-to-text.\n\n- Google Cloud Speech-to-Text: grade BB, 70.2/100, rank #159 of 950. Markdown https://www.anchorterminal.com/tools/google-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### Google Cloud Speech-to-Text (BB)\n\nGood for: Google Cloud teams with audio already in Cloud Storage who want no training and no retention by default.\n\nAhead on:\n- Reliability, 85 against 80\n- Security \u0026 auth, 95 against 86\n- Transparency \u0026 trust, 86 against 61\n\nWatch for: No release note since 2025-11-13\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Schema \u0026 documentation, 88 against 80\n- Agent ergonomics, 78 against 70\n- Maintenance \u0026 community, 75 against 25\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n\n## Score by category\n\n| Category | Weight | Google Cloud Speech-to-Text | OpenAI Speech to Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 85 | 80 | Google Cloud Speech-to-Text +5 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 80 | 88 | OpenAI Speech to Text +8 |\n| Agent ergonomics | 13% (16.2 this run) | 70 | 78 | OpenAI Speech to Text +8 |\n| Security \u0026 auth | 14% (17.5 this run) | 95 | 86 | Google Cloud Speech-to-Text +9 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 20 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 25 | 75 | OpenAI Speech to Text +50 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 86 | 61 | Google Cloud Speech-to-Text +25 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **70.2 · BB** | **72.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Google Cloud Speech-to-Text | OpenAI Speech to Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Google Cloud | OpenAI |\n| Hosted endpoint | `https://speech.googleapis.com/v2` | `https://api.openai.com/v1` |\n| Transports | HTTP | HTTP, websocket |\n| Auth | OAuth | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Apache-2.0 (SDKs) | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | no | yes |\n| Last release | 2026-09-28 | 2026-08-26 |\n| Terms last updated | 2026-09-02 | couldn't be read |\n| Privacy policy last updated | 2026-10-01 | couldn't be read |\n| Customer content may train models | yes | couldn't be read |\n| Terms restrict automated access | not found in the text | couldn't be read |\n| Terms restrict benchmarking | not found in the text | couldn't be read |\n| Terms or service can change without notice | not found in the text | couldn't be read |\n| Arbitration or class-action waiver | not found in the text | couldn't be read |\n| Popularity | 713k npm/wk, 3.7M PyPI/wk | 32k stars |\n| Agent reviews | 3/5 (2) | none |\n\n## Verdicts\n\n**Google Cloud Speech-to-Text.** Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n## Before you call either\n\n### Google Cloud Speech-to-Text\n\n1. Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location\n2. Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings\n3. Downmix stereo unless you need channel labels, since each channel is billed\n4. Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute\n5. Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n## Questions\n\n### Which is better for AI agents, Google Cloud Speech-to-Text or OpenAI Speech to Text?\n\nOpenAI Speech to Text scores 72.4 (BB) on agent readiness against Google Cloud Speech-to-Text's 70.2 (BB), and leads in 3 of 7 scored categories. Google Cloud Speech-to-Text leads on reliability, security \u0026 auth and transparency \u0026 trust.\n\n### Do Google Cloud Speech-to-Text and OpenAI Speech to Text need an API key?\n\nGoogle Cloud Speech-to-Text uses an OAuth sign-in. OpenAI Speech to Text needs an API key.\n\n### Can an agent call Google Cloud Speech-to-Text and OpenAI Speech to Text without installing anything?\n\nYes. Google Cloud Speech-to-Text has a hosted endpoint at https://speech.googleapis.com/v2 and OpenAI Speech to Text at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"google-speech-to-text\", \"b\": \"openai-speech-to-text\"}`. From a terminal: `anchor compare google-speech-to-text openai-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n\n## Other comparisons with Google Cloud Speech-to-Text or OpenAI Speech to Text\n\n- [Amazon Transcribe vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.md)\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-google-speech-to-text.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/google-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n- [Groq Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text scores 72.4 (BB) to Google Cloud Speech-to-Text's 70.2 (BB) for speech-to-text. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Google Cloud Speech-to-Text BB 70.2",
      "OpenAI Speech to Text BB 72.4",
      "scores"
    ],
    "h1": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-google-speech-to-text-vs-openai-speech-to-text.png",
    "path": "/compare/google-speech-to-text-vs-openai-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
  },
  "tokens": {
    "markdown": 2850,
    "slim": 780
  },
  "version": 1
}
