{
  "data": {
    "a": {
      "slug": "assemblyai-stt",
      "name": "AssemblyAI Speech-to-Text (Universal)",
      "vendor": "AssemblyAI",
      "vendorUrl": "https://www.assemblyai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Speech-to-text APIs for recorded audio and live streams, with speaker identification, translation and redaction options.",
      "url": "https://www.anchorterminal.com/tools/assemblyai-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/assemblyai-stt.json",
      "repo": "https://github.com/AssemblyAI/assemblyai-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "streamable-http"
      ],
      "remoteUrl": "https://api.assemblyai.com/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "assemblyai"
        },
        {
          "registry": "pypi",
          "name": "assemblyai"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key in the `authorization` header. Temporary tokens for browser streaming. The docs MCP needs no key.",
      "pricing": "usage",
      "pricingNotes": "Free tier with no card covers up to 185 hours of pre-recorded or 333 hours of streaming. Then pay as you go per hour of audio. Universal-3.5 Pro $0.21, Universal-2 $0.15, Universal-3.6 Pro Realtime $0.45, Universal-Streaming $0.15, Sync $0.45. Diarisation $0.02 an hour on files and $0.12 on streams, keyterms $0.05 on Universal-3.5 Pro, translation $0.06. Streaming bills session time, not audio sent (https://www.assemblyai.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 213,
        "npmWeekly": 599958,
        "pypiWeekly": 738086,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://www.assemblyai.com/docs",
      "llmsTxt": "https://www.assemblyai.com/docs/llms.txt",
      "openapi": "https://www.assemblyai.com/docs/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "no-card",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "mcp",
        "streaming",
        "batch",
        "webhooks",
        "async-jobs",
        "enterprise"
      ],
      "lastRelease": "2026-09-24",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67,
        "grade": "B",
        "agentReady": false,
        "rank": 148,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 8,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 80,
          "maintenance": 80,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 50,
          "transparency": 78
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": -3,
        "negativeNotes": [
          "2026-07-10, the async `speech_model` parameter began returning 400 for current model names and silently routing legacy names to the default model, and `universal-3-pro` was blocked for new accounts and accounts inactive for 7 days, all announced in the changelog the same day. We found no earlier notice (https://www.assemblyai.com/changelog). Documented, so the minimum deduction."
        ],
        "verdict": "OpenAPI 3.1 file with typed inputs and error responses on every operation. Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026.",
        "strengths": [
          "OpenAPI 3.1 file with typed inputs and error responses on every operation",
          "Universal-2 covers 99 languages at $0.15 an hour, Universal-3.5 Pro is $0.21",
          "Transcript lists paginate with `limit`, `before_id` and `after_id`, and transcripts come back as sentences, paragraphs, SRT or VTT",
          "Free tier with no card, up to 185 hours pre-recorded",
          "TTL deletion from 1 hour and a delete endpoint for transcripts"
        ],
        "weaknesses": [
          "Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026",
          "Free accounts can't opt out of model training",
          "The HTTP rate limit returns 403 with no Retry-After",
          "Streaming bills the time the socket is open, up to a 3-hour auto-close",
          "No security.txt or published disclosure policy"
        ],
        "agentNotes": [
          "Send `{\"type\":\"Terminate\"}` to close every stream, or billing runs to the 3-hour auto-close",
          "Use `speech_models` (plural). The singular `speech_model` now returns 400 for current model names",
          "Treat a 403 on polling as the rate limit and back off with jitter, or use webhooks",
          "Fetch `/sentences` or `/paragraphs` instead of the full transcript when you only need text",
          "Opt out in Data Controls on a paid account before sending customer audio"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67
          }
        ],
        "editorialScores": {
          "ergonomics": 80,
          "maintenance": 80,
          "payments": 40,
          "reliability": 70,
          "schema": 95,
          "security": 50,
          "transparency": 70
        },
        "provenanceScore": 86
      },
      "connect": {
        "http": "curl -X POST https://api.assemblyai.com/v2/transcript -H \"authorization: $ASSEMBLYAI_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"audio_url\":\"https://assembly.ai/wildfires.mp3\",\"language_detection\":true,\"speaker_labels\":true}'",
        "claudeCode": "claude mcp add assemblyai-docs --transport http https://assemblyai.com/docs/mcp"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/assemblyai-stt"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Universal-3.5 Pro pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0035,
          "note": "published as $0.21 an hour"
        },
        {
          "item": "Universal-2 pre-recorded",
          "unit": "audio-minute",
          "usd": 0.0025,
          "note": "published as $0.15 an hour"
        },
        {
          "item": "Universal-3.6 Pro Realtime",
          "unit": "audio-minute",
          "usd": 0.0075,
          "note": "published as $0.45 an hour, billed on session time"
        },
        {
          "item": "Universal-Streaming",
          "unit": "audio-minute",
          "usd": 0.0025,
          "note": "published as $0.15 an hour, billed on session time"
        },
        {
          "item": "Sync API",
          "unit": "audio-minute",
          "usd": 0.0075,
          "note": "published as $0.45 an hour, clips up to 2 minutes"
        },
        {
          "item": "Streaming diarisation add-on",
          "unit": "audio-minute",
          "usd": 0.002,
          "note": "published as $0.12 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "AssemblyAI, Inc.",
        "domain": "assemblyai.com",
        "domainRegistered": "2016-12-24",
        "endpointOnVendorDomain": true,
        "terms": "https://www.assemblyai.com/legal/terms-of-service",
        "privacy": "https://www.assemblyai.com/legal/privacy-policy",
        "statusPage": "https://status.assemblyai.com",
        "changelog": "https://www.assemblyai.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 86
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/assemblyai-stt.json",
      "live": {
        "slug": "assemblyai-stt",
        "probe": {
          "target": "https://api.assemblyai.com/v2",
          "method": "get",
          "lastAt": "2026-10-04T23:32:43.1361406Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 444,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 443,
          "p95ms24h": 487,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.assemblyai.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T23:27:39.943165483Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "AssemblyAI/assemblyai-python-sdk",
            "version": "1.6.1",
            "released": "2026-09-24",
            "seenAt": "2026-10-04T16:20:56.490457181Z"
          },
          {
            "registry": "npm",
            "name": "assemblyai",
            "version": "4.41.5",
            "seenAt": "2026-10-04T16:20:55.58059221Z"
          },
          {
            "registry": "pypi",
            "name": "assemblyai",
            "version": "1.6.1",
            "released": "2026-09-24",
            "seenAt": "2026-10-04T16:20:56.382722009Z"
          }
        ],
        "githubStars": 213,
        "npmWeekly": 645780,
        "pypiWeekly": 817313,
        "securityTxt": {
          "url": "https://assemblyai.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:16:02.784112938Z"
        },
        "llmsTxt": {
          "url": "https://www.assemblyai.com/docs/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:15.941338129Z"
        },
        "domain": {
          "domain": "assemblyai.com",
          "registered": "2016-12-24",
          "source": "https://rdap.verisign.com/com/v1/domain/assemblyai.com",
          "checkedAt": "2026-10-04T13:07:48.940216522Z"
        },
        "pages": [
          {
            "url": "https://www.assemblyai.com/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-04T15:49:11.730936349Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "f7dbbe98a56a"
          },
          {
            "url": "https://www.assemblyai.com/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-04T15:49:17.853952095Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "07bb01ccd4d1"
          },
          {
            "url": "https://www.assemblyai.com/legal/privacy-policy",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-04T15:49:13.761556773Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6671be50624e"
          },
          {
            "url": "https://www.assemblyai.com/legal/terms-of-service",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-04T15:49:15.761345368Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "766c6ee47ebc"
          }
        ],
        "updatedAt": "2026-10-04T23:32:43.1361406Z"
      }
    },
    "b": {
      "slug": "google-speech-to-text",
      "name": "Google Cloud Speech-to-Text",
      "vendor": "Google Cloud",
      "vendorUrl": "https://cloud.google.com/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Google Cloud's transcription API.",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "repo": "https://github.com/googleapis/google-cloud-python/tree/main/packages/google-cloud-speech",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://speech.googleapis.com/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-cloud-speech"
        },
        {
          "registry": "npm",
          "name": "@google-cloud/speech"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth 2.0 bearer token from a service account or `gcloud` (Application Default Credentials) on a project with billing and the API turned on. Chirp 3 runs on the `us` and `eu` multi-region endpoints such as `us-speech.googleapis.com`.",
      "pricing": "freemium",
      "pricingNotes": "V2 standard recognition, which covers Chirp 3, is $0.016 a minute to 500,000 minutes a month, then $0.01, $0.008 and $0.004 past 2M. Dynamic batch is $0.003 a minute. Billed per second, per channel. V1 has 60 free minutes a month and charges $0.024 without data logging. Medical models $0.078 (https://cloud.google.com/speech-to-text/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 713013,
        "pypiWeekly": 3703632,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.cloud.google.com/speech-to-text/docs",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "card-required"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 98,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 4,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 88
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
        "strengths": [
          "Audio isn't stored or used for training unless the project opts in to data logging",
          "No Speech-to-Text incident on the Google Cloud status page since 12 June 2025",
          "Dynamic batch at $0.003 a minute, and standard recognition tiers down to $0.004 past 2M minutes",
          "OAuth service accounts with IAM roles, and Cloud Audit Logs",
          "300 concurrent streams per region by default"
        ],
        "weaknesses": [
          "No release note since 2025-11-13",
          "82 of the 111 Chirp 3 locales are preview",
          "Sync requests stop at 1 minute and streams at 5 minutes, and batch reads only from Cloud Storage",
          "No API keys in the documented V2 flow, and the free minutes need a billed project",
          "The quotas page doesn't say what error a breach returns or how to back off"
        ],
        "agentNotes": [
          "Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location",
          "Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings",
          "Downmix stereo unless you need channel labels, since each channel is billed",
          "Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute",
          "Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.4
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 75
        },
        "provenanceScore": 100
      },
      "connect": {
        "install": "pip install google-cloud-speech   # or: npm i @google-cloud/speech",
        "http": "curl -X POST \"https://us-speech.googleapis.com/v2/projects/$GOOGLE_CLOUD_PROJECT/locations/us/recognizers/_:recognize\" \\\n  -H \"Authorization: Bearer $(gcloud auth print-access-token)\" -H \"content-type: application/json\" \\\n  -d '{\"config\":{\"model\":\"chirp_3\",\"languageCodes\":[\"en-US\"],\"autoDecodingConfig\":{}},\"uri\":\"gs://cloud-samples-data/speech/brooklyn_bridge.flac\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/google-speech-to-text"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "google-adk",
        "google-secret-manager",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "google-drive-api",
        "gemini-cli"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "V2 standard recognition (Chirp 3)",
          "unit": "audio-minute",
          "usd": 0.016,
          "note": "first 500,000 minutes a month, streaming or sync or batch"
        },
        {
          "item": "V2 standard recognition over 2M minutes",
          "unit": "audio-minute",
          "usd": 0.004
        },
        {
          "item": "V2 dynamic batch",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "lower-priority batch"
        },
        {
          "item": "V1 without data logging",
          "unit": "audio-minute",
          "usd": 0.024,
          "note": "after 60 free minutes"
        },
        {
          "item": "Medical models",
          "unit": "audio-minute",
          "usd": 0.078
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "domainNote": "The endpoint is on googleapis.com, Google's API domain. google.com was registered in 1997.",
        "endpointOnVendorDomain": true,
        "terms": "https://cloud.google.com/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://status.cloud.google.com",
        "changelog": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 100
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
      "live": {
        "slug": "google-speech-to-text",
        "probe": {
          "target": "https://speech.googleapis.com/v2",
          "method": "get",
          "lastAt": "2026-10-04T23:32:47.953190288Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 31,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 44,
          "p95ms24h": 88,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cloud.google.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-09-30T22:44:37.367865472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "googleapis/google-cloud-python",
            "version": "sqlalchemy-bigquery-v1.17.3",
            "released": "2026-10-02",
            "seenAt": "2026-10-04T16:28:53.455824573Z"
          },
          {
            "registry": "npm",
            "name": "@google-cloud/speech",
            "version": "8.1.1",
            "seenAt": "2026-10-04T16:28:52.627594401Z"
          },
          {
            "registry": "pypi",
            "name": "google-cloud-speech",
            "version": "2.41.0",
            "released": "2026-10-01",
            "seenAt": "2026-10-04T16:28:52.433737787Z"
          }
        ],
        "githubStars": 5400,
        "npmWeekly": 786436,
        "pypiWeekly": 3475501,
        "securityTxt": {
          "url": "https://google.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2030-04-01T00:00:00z",
          "checkedAt": "2026-10-04T15:15:53.387118101Z"
        },
        "domain": {
          "domain": "google.com",
          "registered": "1997-09-15",
          "source": "https://rdap.verisign.com/com/v1/domain/google.com",
          "checkedAt": "2026-10-04T13:05:50.737985829Z"
        },
        "pages": [
          {
            "url": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:29.255732974Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "03f9dac9276b"
          },
          {
            "url": "https://cloud.google.com/speech-to-text/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:41:59.781476607Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "c367763f8641"
          },
          {
            "url": "https://cloud.google.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-01T13:11:34.992628421Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6798e0f4fb24"
          }
        ],
        "updatedAt": "2026-10-04T23:32:47.953190288Z"
      }
    },
    "summary": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against AssemblyAI Speech-to-Text (Universal)'s 67 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.min.md"
  },
  "markdown": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against AssemblyAI Speech-to-Text (Universal)'s 67 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points.\n\n- AssemblyAI Speech-to-Text (Universal): grade B, 67/100, rank #148 of 452. Markdown https://www.anchorterminal.com/tools/assemblyai-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/assemblyai-stt.json\n- Google Cloud Speech-to-Text: grade BB, 70.4/100, rank #98 of 452. Markdown https://www.anchorterminal.com/tools/google-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n\n## Which one, for what\n\nPick AssemblyAI Speech-to-Text (Universal) for schema \u0026 documentation (+15), agent ergonomics (+10), payments \u0026 pricing (+20), maintenance \u0026 community (+55).\n\nPick Google Cloud Speech-to-Text for reliability (+15), security \u0026 auth (+45), transparency \u0026 trust (+10).\n\n## Score by category\n\n| Category | Weight | AssemblyAI Speech-to-Text (Universal) | Google Cloud Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 70 | 85 | Google Cloud Speech-to-Text +15 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 95 | 80 | AssemblyAI Speech-to-Text (Universal) +15 |\n| Agent ergonomics | 13% (16.2 this run) | 80 | 70 | AssemblyAI Speech-to-Text (Universal) +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 50 | 95 | Google Cloud Speech-to-Text +45 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | AssemblyAI Speech-to-Text (Universal) +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 25 | AssemblyAI Speech-to-Text (Universal) +55 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 78 | 88 | Google Cloud Speech-to-Text +10 |\n| Negative events | ≤15 | -3 | 0 | |\n| **Total** | | **67 · B** | **70.4 · BB** | |\n\n## Facts side by side\n\n| Fact | AssemblyAI Speech-to-Text (Universal) | Google Cloud Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | AssemblyAI | Google Cloud |\n| Hosted endpoint | `https://api.assemblyai.com/v2` | `https://speech.googleapis.com/v2` |\n| Transports | HTTP, Streamable HTTP | HTTP |\n| Auth | API key | OAuth |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | MIT (SDKs) | Apache-2.0 (SDKs) |\n| Tools exposed | none | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| MCP registry | not listed | not listed |\n| Last release | 2026-09-24 | 2026-09-28 |\n| Popularity | 213 stars, 600k npm/wk, 738k PyPI/wk | 713k npm/wk, 3.7M PyPI/wk |\n| Agent reviews | 3.5/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**AssemblyAI Speech-to-Text (Universal).** OpenAPI 3.1 file with typed inputs and error responses on every operation. Two outages of an hour or more in the last 90 days, on 31 July and 16 September 2026.\n\n**Google Cloud Speech-to-Text.** Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.\n\n## Before you call either\n\n### AssemblyAI Speech-to-Text (Universal)\n\n1. Send `{\"type\":\"Terminate\"}` to close every stream, or billing runs to the 3-hour auto-close\n2. Use `speech_models` (plural). The singular `speech_model` now returns 400 for current model names\n3. Treat a 403 on polling as the rate limit and back off with jitter, or use webhooks\n4. Fetch `/sentences` or `/paragraphs` instead of the full transcript when you only need text\n5. Opt out in Data Controls on a paid account before sending customer audio\n\n### Google Cloud Speech-to-Text\n\n1. Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location\n2. Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings\n3. Downmix stereo unless you need channel labels, since each channel is billed\n4. Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute\n5. Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval\n\n## Other comparisons with AssemblyAI Speech-to-Text (Universal) or Google Cloud Speech-to-Text\n\n- [Amazon Transcribe vs AssemblyAI Speech-to-Text (Universal)](https://www.anchorterminal.com/compare/amazon-transcribe-vs-assemblyai-stt.md)\n- [Amazon Transcribe vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Azure AI Speech speech-to-text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-azure-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Deepgram Speech-to-Text (Nova-3, Flux)](https://www.anchorterminal.com/compare/assemblyai-stt-vs-deepgram-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs ElevenLabs Scribe Speech to Text API](https://www.anchorterminal.com/compare/assemblyai-stt-vs-elevenlabs-scribe.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/assemblyai-stt-vs-gladia-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/assemblyai-stt-vs-rev-ai-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-soniox-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against AssemblyAI Speech-to-Text (Universal)'s 67 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "AssemblyAI Speech-to-Text (Universal) B 67",
      "Google Cloud Speech-to-Text BB 70.4",
      "scores"
    ],
    "h1": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-assemblyai-stt-vs-google-speech-to-text.png",
    "path": "/compare/assemblyai-stt-vs-google-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text"
  },
  "tokens": {
    "markdown": 1850,
    "slim": 380
  },
  "version": 1
}
