{
  "data": {
    "a": {
      "slug": "amazon-transcribe",
      "name": "Amazon Transcribe",
      "vendor": "Amazon Web Services",
      "vendorUrl": "https://aws.amazon.com/transcribe/",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "AWS's transcription API.",
      "url": "https://www.anchorterminal.com/tools/amazon-transcribe",
      "markdownUrl": "https://www.anchorterminal.com/tools/amazon-transcribe.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/amazon-transcribe.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/amazon-transcribe.json",
      "repo": "https://github.com/awslabs/amazon-transcribe-streaming-sdk",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://transcribe.us-east-1.amazonaws.com",
      "packages": [
        {
          "registry": "pypi",
          "name": "amazon-transcribe"
        },
        {
          "registry": "npm",
          "name": "@aws-sdk/client-transcribe"
        },
        {
          "registry": "npm",
          "name": "@aws-sdk/client-transcribe-streaming"
        }
      ],
      "auth": "api-key",
      "authNotes": "AWS Signature Version 4 with IAM access keys or a role. Batch goes to `transcribe.\u003cregion\u003e.amazonaws.com`, streaming to `transcribestreaming.\u003cregion\u003e.amazonaws.com` (WebSocket needs a presigned URL).",
      "pricing": "usage",
      "pricingNotes": "US East is $0.006 a minute batch and $0.01 a minute streaming, billed per second with no minimum and up to two channels included. Diarisation, custom vocabularies and language ID are included. PII redaction adds $0.0024 a minute and custom language models $0.006. Accounts opened before 2025-07-15 get 60 free minutes a month for 12 months, newer accounts get Free Tier credits instead (https://aws.amazon.com/transcribe/pricing/).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 185,
        "npmWeekly": 505959,
        "pypiWeekly": 200552,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.aws.amazon.com/transcribe/latest/dg/what-is.html",
      "llmsTxt": "https://docs.aws.amazon.com/transcribe/latest/dg/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "hosted",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "llms-txt",
        "card-required"
      ],
      "lastRelease": "2026-09-29",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 73.6,
        "grade": "BB",
        "agentReady": true,
        "rank": 57,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 2,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 80,
          "maintenance": 35,
          "payments": 20,
          "reliability": 95,
          "schema": 90,
          "security": 80,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "$0.006 a minute batch and $0.01 streaming in US East, with diarisation, custom vocabulary and language ID included. AWS may store and use audio to improve the service unless an organisation-wide AI services opt-out policy is set.",
        "strengths": [
          "$0.006 a minute batch and $0.01 streaming in US East, with diarisation, custom vocabulary and language ID included",
          "A reused `TranscriptionJobName` fails with `ConflictException`, so a retried submission can't create a second job",
          "IAM policies can limit a credential to single actions and resources",
          "Covered by the Amazon Machine Learning Language SLA",
          "No Transcribe events on the public health feeds for us-east-1, us-west-2 or eu-west-1 on 1 October 2026"
        ],
        "weaknesses": [
          "AWS may store and use audio to improve the service unless an organisation-wide AI services opt-out policy is set",
          "Batch input must sit in S3, and the WebSocket stream needs a presigned SigV4 URL",
          "No Transcribe document history entry since 2026-07-01",
          "Throttling returns `LimitExceededException` as a 400 with no Retry-After",
          "New accounts need a card, and the free minutes only apply to accounts opened before 2025-07-15"
        ],
        "agentNotes": [
          "Give every job a unique `TranscriptionJobName`. A retry with the same name fails with `ConflictException` rather than starting a second job",
          "Batch is a job. Poll `GetTranscriptionJob` or listen on EventBridge, then fetch the transcript URI",
          "Set `OutputBucketName`, because job records are deleted after 90 days",
          "Back off on `LimitExceededException`. It arrives as a 400, not a 429",
          "Set the organisation's AI services opt-out policy before sending customer audio"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 4,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 73.6
          }
        ],
        "editorialScores": {
          "ergonomics": 80,
          "maintenance": 35,
          "payments": 20,
          "reliability": 95,
          "schema": 90,
          "security": 80,
          "transparency": 75
        },
        "provenanceScore": 95
      },
      "connect": {
        "install": "pip install boto3 amazon-transcribe   # or: npm i @aws-sdk/client-transcribe",
        "http": "curl -X POST \"https://transcribe.us-east-1.amazonaws.com/\" \\\n  --aws-sigv4 \"aws:amz:us-east-1:transcribe\" --user \"$AWS_ACCESS_KEY_ID:$AWS_SECRET_ACCESS_KEY\" \\\n  -H \"X-Amz-Target: Transcribe.StartTranscriptionJob\" -H \"content-type: application/x-amz-json-1.1\" \\\n  -d '{\"TranscriptionJobName\":\"call-001\",\"LanguageCode\":\"en-US\",\"Media\":{\"MediaFileUri\":\"s3://my-bucket/call.wav\"}}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/amazon-transcribe"
      },
      "sameCompany": [
        "amazon-bedrock-guardrails",
        "amazon-polly",
        "aws-secrets-manager",
        "aws-mcp-servers",
        "amazon-ses",
        "amazon-translate"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Batch",
          "unit": "audio-minute",
          "usd": 0.006,
          "note": "US East, per second"
        },
        {
          "item": "Streaming",
          "unit": "audio-minute",
          "usd": 0.01,
          "note": "US East, per second"
        },
        {
          "item": "PII redaction add-on",
          "unit": "audio-minute",
          "usd": 0.0024,
          "note": "first 250,000 minutes"
        },
        {
          "item": "Custom language model add-on",
          "unit": "audio-minute",
          "usd": 0.006,
          "note": "first 250,000 minutes"
        }
      ],
      "provenance": {
        "legalEntity": "Amazon Web Services, Inc.",
        "domain": "amazon.com",
        "domainRegistered": "1994-11-01",
        "domainNote": "The endpoints are on amazonaws.com (registered 2005-08-18) and api.aws, both AWS domains. The security.txt on aws.amazon.com passed its Expires date on 2026-09-24.",
        "endpointOnVendorDomain": true,
        "terms": "https://aws.amazon.com/service-terms/",
        "privacy": "https://aws.amazon.com/privacy/",
        "statusPage": "https://health.aws.amazon.com/health/status",
        "changelog": "https://docs.aws.amazon.com/transcribe/latest/dg/doc-history.html",
        "securityTxt": "expired",
        "checked": "2026-09-30",
        "score": 95
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/amazon-transcribe.json",
      "live": {
        "slug": "amazon-transcribe",
        "probe": {
          "target": "https://transcribe.us-east-1.amazonaws.com",
          "method": "get",
          "lastAt": "2026-10-04T23:32:42.85803071Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 263,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 255,
          "p95ms24h": 314,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@aws-sdk/client-transcribe",
            "version": "3.1146.0",
            "seenAt": "2026-10-04T16:20:14.407745351Z"
          },
          {
            "registry": "npm",
            "name": "@aws-sdk/client-transcribe-streaming",
            "version": "3.1146.0",
            "seenAt": "2026-10-04T16:20:15.237150076Z"
          },
          {
            "registry": "pypi",
            "name": "amazon-transcribe",
            "version": "0.6.4",
            "released": "2025-05-05",
            "seenAt": "2026-10-04T16:20:14.218395026Z"
          }
        ],
        "githubStars": 185,
        "npmWeekly": 609911,
        "pypiWeekly": 233186,
        "securityTxt": {
          "url": "https://amazon.com/.well-known/security.txt",
          "state": "valid",
          "checkedAt": "2026-10-04T15:15:49.458098289Z"
        },
        "llmsTxt": {
          "url": "https://docs.aws.amazon.com/transcribe/latest/dg/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:18.885309232Z"
        },
        "domain": {
          "domain": "amazon.com",
          "registered": "1994-11-01",
          "source": "https://rdap.verisign.com/com/v1/domain/amazon.com",
          "checkedAt": "2026-10-04T13:06:18.739682554Z"
        },
        "pages": [
          {
            "url": "https://docs.aws.amazon.com/transcribe/latest/dg/doc-history.html",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-04T15:43:22.418825271Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "2b53ad2d121d"
          },
          {
            "url": "https://aws.amazon.com/transcribe/pricing/",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-04T15:41:40.655188059Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "9d86c82cdd96"
          }
        ],
        "updatedAt": "2026-10-04T23:32:42.85803071Z"
      }
    },
    "b": {
      "slug": "google-speech-to-text",
      "name": "Google Cloud Speech-to-Text",
      "vendor": "Google Cloud",
      "vendorUrl": "https://cloud.google.com/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Google Cloud's transcription API.",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "repo": "https://github.com/googleapis/google-cloud-python/tree/main/packages/google-cloud-speech",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://speech.googleapis.com/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-cloud-speech"
        },
        {
          "registry": "npm",
          "name": "@google-cloud/speech"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth 2.0 bearer token from a service account or `gcloud` (Application Default Credentials) on a project with billing and the API turned on. Chirp 3 runs on the `us` and `eu` multi-region endpoints such as `us-speech.googleapis.com`.",
      "pricing": "freemium",
      "pricingNotes": "V2 standard recognition, which covers Chirp 3, is $0.016 a minute to 500,000 minutes a month, then $0.01, $0.008 and $0.004 past 2M. Dynamic batch is $0.003 a minute. Billed per second, per channel. V1 has 60 free minutes a month and charges $0.024 without data logging. Medical models $0.078 (https://cloud.google.com/speech-to-text/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 713013,
        "pypiWeekly": 3703632,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.cloud.google.com/speech-to-text/docs",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "card-required"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 98,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 4,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 88
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
        "strengths": [
          "Audio isn't stored or used for training unless the project opts in to data logging",
          "No Speech-to-Text incident on the Google Cloud status page since 12 June 2025",
          "Dynamic batch at $0.003 a minute, and standard recognition tiers down to $0.004 past 2M minutes",
          "OAuth service accounts with IAM roles, and Cloud Audit Logs",
          "300 concurrent streams per region by default"
        ],
        "weaknesses": [
          "No release note since 2025-11-13",
          "82 of the 111 Chirp 3 locales are preview",
          "Sync requests stop at 1 minute and streams at 5 minutes, and batch reads only from Cloud Storage",
          "No API keys in the documented V2 flow, and the free minutes need a billed project",
          "The quotas page doesn't say what error a breach returns or how to back off"
        ],
        "agentNotes": [
          "Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location",
          "Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings",
          "Downmix stereo unless you need channel labels, since each channel is billed",
          "Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute",
          "Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.4
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 75
        },
        "provenanceScore": 100
      },
      "connect": {
        "install": "pip install google-cloud-speech   # or: npm i @google-cloud/speech",
        "http": "curl -X POST \"https://us-speech.googleapis.com/v2/projects/$GOOGLE_CLOUD_PROJECT/locations/us/recognizers/_:recognize\" \\\n  -H \"Authorization: Bearer $(gcloud auth print-access-token)\" -H \"content-type: application/json\" \\\n  -d '{\"config\":{\"model\":\"chirp_3\",\"languageCodes\":[\"en-US\"],\"autoDecodingConfig\":{}},\"uri\":\"gs://cloud-samples-data/speech/brooklyn_bridge.flac\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/google-speech-to-text"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "google-adk",
        "google-secret-manager",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "google-drive-api",
        "gemini-cli"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "V2 standard recognition (Chirp 3)",
          "unit": "audio-minute",
          "usd": 0.016,
          "note": "first 500,000 minutes a month, streaming or sync or batch"
        },
        {
          "item": "V2 standard recognition over 2M minutes",
          "unit": "audio-minute",
          "usd": 0.004
        },
        {
          "item": "V2 dynamic batch",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "lower-priority batch"
        },
        {
          "item": "V1 without data logging",
          "unit": "audio-minute",
          "usd": 0.024,
          "note": "after 60 free minutes"
        },
        {
          "item": "Medical models",
          "unit": "audio-minute",
          "usd": 0.078
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "domainNote": "The endpoint is on googleapis.com, Google's API domain. google.com was registered in 1997.",
        "endpointOnVendorDomain": true,
        "terms": "https://cloud.google.com/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://status.cloud.google.com",
        "changelog": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 100
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
      "live": {
        "slug": "google-speech-to-text",
        "probe": {
          "target": "https://speech.googleapis.com/v2",
          "method": "get",
          "lastAt": "2026-10-04T23:32:47.953190288Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 31,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 44,
          "p95ms24h": 88,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cloud.google.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-09-30T22:44:37.367865472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "googleapis/google-cloud-python",
            "version": "sqlalchemy-bigquery-v1.17.3",
            "released": "2026-10-02",
            "seenAt": "2026-10-04T16:28:53.455824573Z"
          },
          {
            "registry": "npm",
            "name": "@google-cloud/speech",
            "version": "8.1.1",
            "seenAt": "2026-10-04T16:28:52.627594401Z"
          },
          {
            "registry": "pypi",
            "name": "google-cloud-speech",
            "version": "2.41.0",
            "released": "2026-10-01",
            "seenAt": "2026-10-04T16:28:52.433737787Z"
          }
        ],
        "githubStars": 5400,
        "npmWeekly": 786436,
        "pypiWeekly": 3475501,
        "securityTxt": {
          "url": "https://google.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2030-04-01T00:00:00z",
          "checkedAt": "2026-10-04T15:15:53.387118101Z"
        },
        "domain": {
          "domain": "google.com",
          "registered": "1997-09-15",
          "source": "https://rdap.verisign.com/com/v1/domain/google.com",
          "checkedAt": "2026-10-04T13:05:50.737985829Z"
        },
        "pages": [
          {
            "url": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:29.255732974Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "03f9dac9276b"
          },
          {
            "url": "https://cloud.google.com/speech-to-text/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:41:59.781476607Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "c367763f8641"
          },
          {
            "url": "https://cloud.google.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-01T13:11:34.992628421Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6798e0f4fb24"
          }
        ],
        "updatedAt": "2026-10-04T23:32:47.953190288Z"
      }
    },
    "summary": "Amazon Transcribe has a score of 73.6 (BB) against Google Cloud Speech-to-Text's 70.4 (BB). Both do speech stt. The largest gap is security \u0026 auth, 15 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.min.md"
  },
  "markdown": "Amazon Transcribe has a score of 73.6 (BB) against Google Cloud Speech-to-Text's 70.4 (BB). Both do speech stt. The largest gap is security \u0026 auth, 15 points.\n\n- Amazon Transcribe: grade BB, 73.6/100, rank #57 of 452. Markdown https://www.anchorterminal.com/tools/amazon-transcribe.md · JSON https://www.anchorterminal.com/api/v1/tools/amazon-transcribe.json\n- Google Cloud Speech-to-Text: grade BB, 70.4/100, rank #98 of 452. Markdown https://www.anchorterminal.com/tools/google-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n\n## Which one, for what\n\nPick Amazon Transcribe for reliability (+10), schema \u0026 documentation (+10), agent ergonomics (+10), maintenance \u0026 community (+10).\n\nPick Google Cloud Speech-to-Text for security \u0026 auth (+15).\n\n## Score by category\n\n| Category | Weight | Amazon Transcribe | Google Cloud Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 95 | 85 | Amazon Transcribe +10 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 90 | 80 | Amazon Transcribe +10 |\n| Agent ergonomics | 13% (16.2 this run) | 80 | 70 | Amazon Transcribe +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 80 | 95 | Google Cloud Speech-to-Text +15 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 20 | 20 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 35 | 25 | Amazon Transcribe +10 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 85 | 88 | Google Cloud Speech-to-Text +3 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **73.6 · BB** | **70.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Amazon Transcribe | Google Cloud Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Amazon Web Services | Google Cloud |\n| Hosted endpoint | `https://transcribe.us-east-1.amazonaws.com` | `https://speech.googleapis.com/v2` |\n| Transports | HTTP | HTTP |\n| Auth | API key | OAuth |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | Apache-2.0 (SDKs) | Apache-2.0 (SDKs) |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| MCP registry | not listed | not listed |\n| Last release | 2026-09-29 | 2026-09-28 |\n| Popularity | 185 stars, 506k npm/wk, 201k PyPI/wk | 713k npm/wk, 3.7M PyPI/wk |\n| Agent reviews | 4/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**Amazon Transcribe.** $0.006 a minute batch and $0.01 streaming in US East, with diarisation, custom vocabulary and language ID included. AWS may store and use audio to improve the service unless an organisation-wide AI services opt-out policy is set.\n\n**Google Cloud Speech-to-Text.** Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.\n\n## Before you call either\n\n### Amazon Transcribe\n\n1. Give every job a unique `TranscriptionJobName`. A retry with the same name fails with `ConflictException` rather than starting a second job\n2. Batch is a job. Poll `GetTranscriptionJob` or listen on EventBridge, then fetch the transcript URI\n3. Set `OutputBucketName`, because job records are deleted after 90 days\n4. Back off on `LimitExceededException`. It arrives as a 400, not a 429\n5. Set the organisation's AI services opt-out policy before sending customer audio\n\n### Google Cloud Speech-to-Text\n\n1. Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location\n2. Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings\n3. Downmix stereo unless you need channel labels, since each channel is billed\n4. Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute\n5. Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval\n\n## Other comparisons with Amazon Transcribe or Google Cloud Speech-to-Text\n\n- [Amazon Transcribe vs AssemblyAI Speech-to-Text (Universal)](https://www.anchorterminal.com/compare/amazon-transcribe-vs-assemblyai-stt.md)\n- [Amazon Transcribe vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-azure-speech-to-text.md)\n- [Amazon Transcribe vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/amazon-transcribe-vs-deepgram-stt.md)\n- [Amazon Transcribe vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/amazon-transcribe-vs-elevenlabs-scribe.md)\n- [Amazon Transcribe vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/amazon-transcribe-vs-gladia-stt.md)\n- [Amazon Transcribe vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/amazon-transcribe-vs-rev-ai-stt.md)\n- [Amazon Transcribe vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-soniox-stt.md)\n- [Amazon Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Amazon Transcribe vs Google Cloud Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Amazon Transcribe has a score of 73.6 (BB) against Google Cloud Speech-to-Text's 70.4 (BB). Both do speech stt. The largest gap is security \u0026 auth, 15 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Amazon Transcribe BB 73.6",
      "Google Cloud Speech-to-Text BB 70.4",
      "scores"
    ],
    "h1": "Amazon Transcribe vs Google Cloud Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-amazon-transcribe-vs-google-speech-to-text.png",
    "path": "/compare/amazon-transcribe-vs-google-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Amazon Transcribe vs Google Cloud Speech-to-Text for AI agents",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text"
  },
  "tokens": {
    "markdown": 1750,
    "slim": 380
  },
  "version": 1
}
