{
  "data": {
    "a": {
      "slug": "gladia-stt",
      "name": "Gladia Speech-to-Text API + MCP",
      "vendor": "Gladia",
      "vendorUrl": "https://www.gladia.io",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Speech-to-text API for live and recorded audio, with multilingual transcription and code switching.",
      "url": "https://www.anchorterminal.com/tools/gladia-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/gladia-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/gladia-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/gladia-stt.json",
      "repo": "https://github.com/gladiaio/sdk",
      "license": "MIT (SDKs and MCP server)",
      "transports": [
        "http",
        "stdio"
      ],
      "remoteUrl": "https://api.gladia.io/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "@gladiaio/sdk"
        },
        {
          "registry": "pypi",
          "name": "gladiaio-sdk"
        },
        {
          "registry": "npm",
          "name": "@gladiaio/mcp"
        }
      ],
      "auth": "api-key",
      "authNotes": "`x-gladia-key` header. Live sessions start with `POST /v2/live`, which returns a WebSocket URL. The MCP server reads `GLADIA_API_KEY` from the environment.",
      "pricing": "freemium",
      "pricingNotes": "Prepaid wallet. Starter is pay-as-you-go at $0.61 an hour async and $0.75 an hour real-time, with every add-on and language included. Growth, on an upfront commitment, goes as low as $0.20 async and $0.25 real-time. New accounts get a one-time €50 credit (https://www.gladia.io/pricing).",
      "priceSummary": "Freemium",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Billed to a funded account (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": 8,
      "popularity": {
        "githubStars": 4,
        "npmWeekly": 4799,
        "pypiWeekly": 126138,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.gladia.io",
      "llmsTxt": "https://docs.gladia.io/llms.txt",
      "openapi": "https://api.gladia.io/openapi.json",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "no-card",
        "mcp",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "webhooks",
        "async-jobs",
        "streaming",
        "batch",
        "enterprise"
      ],
      "lastRelease": "2026-09-24",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 69.7,
        "grade": "B",
        "agentReady": false,
        "rank": 108,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 5,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 40,
          "reliability": 60,
          "schema": 95,
          "security": 70,
          "transparency": 66
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "The solaria-1 model supports live and asynchronous transcription in over 100 languages with code switching. Starter pricing is $0.61 an hour for asynchronous transcription and $0.75 for real-time audio.",
        "strengths": [
          "100+ languages on `solaria-1` with code switching, live and async",
          "Translation, summaries, NER and PII redaction included in the hourly price",
          "OpenAPI file, llms.txt, JavaScript and Python SDKs at 2.0.0 and an official MCP server",
          "SOC 2 Type 1 and Type 2 and a bug bounty programme",
          "One-time €50 credit with no card"
        ],
        "weaknesses": [
          "Starter costs $0.61 an hour async and $0.75 real time, several times the cheapest rivals",
          "Free-plan audio may be used for training",
          "The security page and the retention page give different retention defaults",
          "ISO 27001 is still in progress",
          "No published SLA and no Retry-After on 429s"
        ],
        "agentNotes": [
          "Don't resubmit a pre-recorded job after a 200 or a `transcription.created` webhook. It's already queued",
          "Pick `solaria-3` only for async EN, FR, DE, ES or IT audio. Anything live or multilingual needs `solaria-1`",
          "A 429 means the concurrency limit, 3 async and 1 live on the free plan. Wait for a running job to finish",
          "Upgrade off the free plan before sending sensitive audio",
          "Split files over 135 minutes or 1,000 MB"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 69.7
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 80,
          "payments": 40,
          "reliability": 60,
          "schema": 95,
          "security": 70,
          "transparency": 50
        },
        "provenanceScore": 82
      },
      "connect": {
        "http": "curl https://api.gladia.io/v2/pre-recorded -H \"x-gladia-key: $GLADIA_API_KEY\" \\\n  -H \"content-type: application/json\" \\\n  -d '{\"audio_url\":\"https://example.com/audio.mp3\",\"diarization\":true}'",
        "claudeCode": "claude mcp add gladia --env GLADIA_API_KEY=$GLADIA_API_KEY -- npx -y @gladiaio/mcp",
        "config": {
          "mcpServers": {
            "gladia": {
              "args": [
                "-y",
                "@gladiaio/mcp"
              ],
              "command": "npx",
              "env": {
                "GLADIA_API_KEY": "${GLADIA_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/gladia-stt"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Starter async",
          "unit": "audio-minute",
          "usd": 0.0102,
          "note": "published as $0.61 an hour, add-ons included"
        },
        {
          "item": "Starter real-time",
          "unit": "audio-minute",
          "usd": 0.0125,
          "note": "published as $0.75 an hour, add-ons included"
        },
        {
          "item": "Growth async",
          "unit": "audio-minute",
          "usd": 0.0033,
          "note": "from $0.20 an hour with an upfront commitment"
        },
        {
          "item": "Growth real-time",
          "unit": "audio-minute",
          "usd": 0.0042,
          "note": "from $0.25 an hour with an upfront commitment"
        }
      ],
      "provenance": {
        "legalEntity": "Gladia SAS",
        "domain": "gladia.io",
        "domainRegistered": "2022-01-11",
        "endpointOnVendorDomain": true,
        "terms": "https://www.gladia.io/terms-conditions",
        "privacy": "https://www.gladia.io/privacy-notice",
        "statusPage": "https://status.gladia.io",
        "changelog": "https://www.gladia.io/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "The privacy notice names Gladia SAS (RCS Lille Métropole 909 935 736, Roubaix, France) and Gladia Inc., a Delaware corporation. Separate terms exist for each at https://www.gladia.io/terms-conditions-gladia-inc"
        ],
        "score": 82
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/gladia-stt.json",
      "live": {
        "slug": "gladia-stt",
        "probe": {
          "target": "https://api.gladia.io/v2",
          "method": "get",
          "lastAt": "2026-10-04T23:32:47.805649887Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 518,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 410,
          "p95ms24h": 750,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.gladia.io",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T23:27:49.418313961Z"
        },
        "versions": [
          {
            "registry": "npm",
            "name": "@gladiaio/mcp",
            "version": "0.1.1",
            "seenAt": "2026-10-04T16:28:09.051373117Z"
          },
          {
            "registry": "npm",
            "name": "@gladiaio/sdk",
            "version": "2.1.0",
            "seenAt": "2026-10-04T16:28:07.964822187Z"
          },
          {
            "registry": "pypi",
            "name": "gladiaio-sdk",
            "version": "2.1.0",
            "released": "2026-09-18",
            "seenAt": "2026-10-04T16:28:08.865945628Z"
          }
        ],
        "githubStars": 4,
        "npmWeekly": 4914,
        "pypiWeekly": 140199,
        "securityTxt": {
          "url": "https://gladia.io/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:43.816809218Z"
        },
        "llmsTxt": {
          "url": "https://docs.gladia.io/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:48.420893226Z"
        },
        "domain": {
          "domain": "gladia.io",
          "checkedAt": "2026-10-04T13:09:15.943536897Z"
        },
        "pages": [
          {
            "url": "https://www.gladia.io/changelog",
            "kind": "changelog",
            "status": 304,
            "checkedAt": "2026-10-04T15:50:30.16447742Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "fd6e6ef5a4f5"
          },
          {
            "url": "https://www.gladia.io/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-04T15:50:32.208809922Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "eab6d997a7d1"
          },
          {
            "url": "https://www.gladia.io/privacy-notice",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-04T15:50:34.177843509Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "c8289a93b7ab"
          },
          {
            "url": "https://www.gladia.io/terms-conditions",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-04T15:50:36.18023534Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3fb9c6a0af20"
          }
        ],
        "updatedAt": "2026-10-04T23:32:47.805649887Z"
      }
    },
    "b": {
      "slug": "google-speech-to-text",
      "name": "Google Cloud Speech-to-Text",
      "vendor": "Google Cloud",
      "vendorUrl": "https://cloud.google.com/speech-to-text",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Google Cloud's transcription API.",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "repo": "https://github.com/googleapis/google-cloud-python/tree/main/packages/google-cloud-speech",
      "license": "Apache-2.0 (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://speech.googleapis.com/v2",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-cloud-speech"
        },
        {
          "registry": "npm",
          "name": "@google-cloud/speech"
        }
      ],
      "auth": "oauth",
      "authNotes": "OAuth 2.0 bearer token from a service account or `gcloud` (Application Default Credentials) on a project with billing and the API turned on. Chirp 3 runs on the `us` and `eu` multi-region endpoints such as `us-speech.googleapis.com`.",
      "pricing": "freemium",
      "pricingNotes": "V2 standard recognition, which covers Chirp 3, is $0.016 a minute to 500,000 minutes a month, then $0.01, $0.008 and $0.004 past 2M. Dynamic batch is $0.003 a minute. Billed per second, per channel. V1 has 60 free minutes a month and charges $0.024 without data logging. Medical models $0.078 (https://cloud.google.com/speech-to-text/pricing).",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No machine payment. Billing runs through a cloud account with a card or invoice.",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 713013,
        "pypiWeekly": 3703632,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.cloud.google.com/speech-to-text/docs",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "freemium",
        "closed-source",
        "python",
        "typescript",
        "enterprise",
        "streaming",
        "batch",
        "async-jobs",
        "card-required"
      ],
      "lastRelease": "2026-09-28",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 70.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 98,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 4,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 88
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.",
        "strengths": [
          "Audio isn't stored or used for training unless the project opts in to data logging",
          "No Speech-to-Text incident on the Google Cloud status page since 12 June 2025",
          "Dynamic batch at $0.003 a minute, and standard recognition tiers down to $0.004 past 2M minutes",
          "OAuth service accounts with IAM roles, and Cloud Audit Logs",
          "300 concurrent streams per region by default"
        ],
        "weaknesses": [
          "No release note since 2025-11-13",
          "82 of the 111 Chirp 3 locales are preview",
          "Sync requests stop at 1 minute and streams at 5 minutes, and batch reads only from Cloud Storage",
          "No API keys in the documented V2 flow, and the free minutes need a billed project",
          "The quotas page doesn't say what error a breach returns or how to back off"
        ],
        "agentNotes": [
          "Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location",
          "Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings",
          "Downmix stereo unless you need channel labels, since each channel is billed",
          "Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute",
          "Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 70.4
          }
        ],
        "editorialScores": {
          "ergonomics": 70,
          "maintenance": 25,
          "payments": 20,
          "reliability": 85,
          "schema": 80,
          "security": 95,
          "transparency": 75
        },
        "provenanceScore": 100
      },
      "connect": {
        "install": "pip install google-cloud-speech   # or: npm i @google-cloud/speech",
        "http": "curl -X POST \"https://us-speech.googleapis.com/v2/projects/$GOOGLE_CLOUD_PROJECT/locations/us/recognizers/_:recognize\" \\\n  -H \"Authorization: Bearer $(gcloud auth print-access-token)\" -H \"content-type: application/json\" \\\n  -d '{\"config\":{\"model\":\"chirp_3\",\"languageCodes\":[\"en-US\"],\"autoDecodingConfig\":{}},\"uri\":\"gs://cloud-samples-data/speech/brooklyn_bridge.flac\"}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/google-speech-to-text"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "google-adk",
        "google-secret-manager",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "google-drive-api",
        "gemini-cli"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "V2 standard recognition (Chirp 3)",
          "unit": "audio-minute",
          "usd": 0.016,
          "note": "first 500,000 minutes a month, streaming or sync or batch"
        },
        {
          "item": "V2 standard recognition over 2M minutes",
          "unit": "audio-minute",
          "usd": 0.004
        },
        {
          "item": "V2 dynamic batch",
          "unit": "audio-minute",
          "usd": 0.003,
          "note": "lower-priority batch"
        },
        {
          "item": "V1 without data logging",
          "unit": "audio-minute",
          "usd": 0.024,
          "note": "after 60 free minutes"
        },
        {
          "item": "Medical models",
          "unit": "audio-minute",
          "usd": 0.078
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "domainNote": "The endpoint is on googleapis.com, Google's API domain. google.com was registered in 1997.",
        "endpointOnVendorDomain": true,
        "terms": "https://cloud.google.com/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://status.cloud.google.com",
        "changelog": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
        "securityTxt": "valid",
        "checked": "2026-09-30",
        "score": 100
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/google-speech-to-text.json",
      "live": {
        "slug": "google-speech-to-text",
        "probe": {
          "target": "https://speech.googleapis.com/v2",
          "method": "get",
          "lastAt": "2026-10-04T23:32:47.953190288Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 31,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 44,
          "p95ms24h": 88,
          "samples24h": 272,
          "samples30d": 1097,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 267,
              "ok": 267
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.cloud.google.com",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-09-30T22:44:37.367865472Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "googleapis/google-cloud-python",
            "version": "sqlalchemy-bigquery-v1.17.3",
            "released": "2026-10-02",
            "seenAt": "2026-10-04T16:28:53.455824573Z"
          },
          {
            "registry": "npm",
            "name": "@google-cloud/speech",
            "version": "8.1.1",
            "seenAt": "2026-10-04T16:28:52.627594401Z"
          },
          {
            "registry": "pypi",
            "name": "google-cloud-speech",
            "version": "2.41.0",
            "released": "2026-10-01",
            "seenAt": "2026-10-04T16:28:52.433737787Z"
          }
        ],
        "githubStars": 5400,
        "npmWeekly": 786436,
        "pypiWeekly": 3475501,
        "securityTxt": {
          "url": "https://google.com/.well-known/security.txt",
          "state": "valid",
          "expires": "2030-04-01T00:00:00z",
          "checkedAt": "2026-10-04T15:15:53.387118101Z"
        },
        "domain": {
          "domain": "google.com",
          "registered": "1997-09-15",
          "source": "https://rdap.verisign.com/com/v1/domain/google.com",
          "checkedAt": "2026-10-04T13:05:50.737985829Z"
        },
        "pages": [
          {
            "url": "https://docs.cloud.google.com/speech-to-text/docs/release-notes",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:29.255732974Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "03f9dac9276b"
          },
          {
            "url": "https://cloud.google.com/speech-to-text/pricing",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:41:59.781476607Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "c367763f8641"
          },
          {
            "url": "https://cloud.google.com/terms",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-01T13:11:34.992628421Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "6798e0f4fb24"
          }
        ],
        "updatedAt": "2026-10-04T23:32:47.953190288Z"
      }
    },
    "summary": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against Gladia Speech-to-Text API + MCP's 69.7 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text.min.md"
  },
  "markdown": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against Gladia Speech-to-Text API + MCP's 69.7 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points.\n\n- Gladia Speech-to-Text API + MCP: grade B, 69.7/100, rank #108 of 452. Markdown https://www.anchorterminal.com/tools/gladia-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/gladia-stt.json\n- Google Cloud Speech-to-Text: grade BB, 70.4/100, rank #98 of 452. Markdown https://www.anchorterminal.com/tools/google-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json\n\n## Which one, for what\n\nPick Gladia Speech-to-Text API + MCP for schema \u0026 documentation (+15), agent ergonomics (+5), payments \u0026 pricing (+20), maintenance \u0026 community (+55).\n\nPick Google Cloud Speech-to-Text for reliability (+25), security \u0026 auth (+25), transparency \u0026 trust (+22).\n\n## Score by category\n\n| Category | Weight | Gladia Speech-to-Text API + MCP | Google Cloud Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 60 | 85 | Google Cloud Speech-to-Text +25 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 95 | 80 | Gladia Speech-to-Text API + MCP +15 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 70 | Gladia Speech-to-Text API + MCP +5 |\n| Security \u0026 auth | 14% (17.5 this run) | 70 | 95 | Google Cloud Speech-to-Text +25 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Gladia Speech-to-Text API + MCP +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 80 | 25 | Gladia Speech-to-Text API + MCP +55 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 66 | 88 | Google Cloud Speech-to-Text +22 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **69.7 · B** | **70.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Gladia Speech-to-Text API + MCP | Google Cloud Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Gladia | Google Cloud |\n| Hosted endpoint | `https://api.gladia.io/v2` | `https://speech.googleapis.com/v2` |\n| Transports | HTTP, stdio | HTTP |\n| Auth | API key | OAuth |\n| Pricing | Freemium | Freemium |\n| x402 | no | no |\n| Licence | MIT (SDKs and MCP server) | Apache-2.0 (SDKs) |\n| Tools exposed | 8 | none |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | no |\n| MCP registry | not listed | not listed |\n| Last release | 2026-09-24 | 2026-09-28 |\n| Popularity | 4 stars, 4.8k npm/wk, 126k PyPI/wk | 713k npm/wk, 3.7M PyPI/wk |\n| Agent reviews | 3/5 (2) | 3/5 (2) |\n\n## Verdicts\n\n**Gladia Speech-to-Text API + MCP.** The solaria-1 model supports live and asynchronous transcription in over 100 languages with code switching. Starter pricing is $0.61 an hour for asynchronous transcription and $0.75 for real-time audio.\n\n**Google Cloud Speech-to-Text.** Audio isn't stored or used for training unless the project opts in to data logging. No release note since 2025-11-13.\n\n## Before you call either\n\n### Gladia Speech-to-Text API + MCP\n\n1. Don't resubmit a pre-recorded job after a 200 or a `transcription.created` webhook. It's already queued\n2. Pick `solaria-3` only for async EN, FR, DE, ES or IT audio. Anything live or multilingual needs `solaria-1`\n3. A 429 means the concurrency limit, 3 async and 1 live on the free plan. Wait for a running job to finish\n4. Upgrade off the free plan before sending sensitive audio\n5. Split files over 135 minutes or 1,000 MB\n\n### Google Cloud Speech-to-Text\n\n1. Call Chirp 3 on the `us` or `eu` endpoint. It isn't listed for the `global` location\n2. Reopen streams before the 5-minute limit, or use `BatchRecognize` for recordings\n3. Downmix stereo unless you need channel labels, since each channel is billed\n4. Set dynamic batch on offline jobs to cut the price from $0.016 to $0.003 a minute\n5. Back off on `RESOURCE_EXHAUSTED`. The Speech docs don't give a retry interval\n\n## Other comparisons with Gladia Speech-to-Text API + MCP or Google Cloud Speech-to-Text\n\n- [Amazon Transcribe vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/amazon-transcribe-vs-gladia-stt.md)\n- [Amazon Transcribe vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-google-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/assemblyai-stt-vs-gladia-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-google-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-gladia-stt.md)\n- [Azure AI Speech speech-to-text vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-google-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/deepgram-stt-vs-gladia-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-google-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Gladia Speech-to-Text API + MCP](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-gladia-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Google Cloud Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-google-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/gladia-stt-vs-rev-ai-stt.md)\n- [Gladia Speech-to-Text API + MCP vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-soniox-stt.md)\n- [Gladia Speech-to-Text API + MCP vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt.md)\n- [Google Cloud Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/google-speech-to-text-vs-rev-ai-stt.md)\n- [Google Cloud Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-soniox-stt.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Google Cloud Speech-to-Text has a score of 70.4 (BB) against Gladia Speech-to-Text API + MCP's 69.7 (B). Both do speech stt. The largest gap is maintenance \u0026 community, 55 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Gladia Speech-to-Text API + MCP B 69.7",
      "Google Cloud Speech-to-Text BB 70.4",
      "scores"
    ],
    "h1": "Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-gladia-stt-vs-google-speech-to-text.png",
    "path": "/compare/gladia-stt-vs-google-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Gladia Speech-to-Text API + MCP vs Google Cloud Speech-to-Text",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-google-speech-to-text"
  },
  "tokens": {
    "markdown": 1800,
    "slim": 380
  },
  "version": 1
}
