{
  "data": {
    "a": {
      "slug": "groq-speech-to-text",
      "name": "Groq Speech-to-Text",
      "vendor": "Groq",
      "vendorUrl": "https://groq.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Groq's hosted speech-to-text API. It runs OpenAI's Whisper Large v3 and Whisper Large v3 Turbo on OpenAI-compatible transcription and translation endpoints, for uploaded files or audio URLs, with a half-price batch mode.",
      "url": "https://www.anchorterminal.com/tools/groq-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json",
      "repo": "https://github.com/groq/groq-python",
      "license": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.groq.com/openai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "groq"
        },
        {
          "registry": "npm",
          "name": "groq-sdk"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from the GroqCloud console, created inside a project. Projects carry their own rate limits per model, usage data and request logs (https://console.groq.com/docs/projects).",
      "pricing": "freemium",
      "pricingNotes": "$0.04 per audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with a 10-second minimum a request and 50% off through the Batch API (https://console.groq.com/docs/models). The free plan needs no card, so an agent's owner can start without a contract. The Developer plan is postpaid by card, US bank account or SEPA debit.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the speech-to-text guide, the API reference or the billing pages (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 621,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://console.groq.com/docs/speech-to-text",
      "llmsTxt": "https://console.groq.com/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "fast",
        "free-tier",
        "no-card",
        "llms-txt",
        "python",
        "typescript",
        "batch",
        "openai-compatible"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 71.8,
        "grade": "BB",
        "agentReady": true,
        "rank": 113,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.",
        "bestFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "strengths": [
          "Published prices of $0.04 an audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with 50% off through the Batch API",
          "Free plan with no card at 20 requests a minute, 2,000 a day and 7,200 audio seconds an hour on both models",
          "Inputs and outputs are not retained by default, and zero data retention is a console setting that covers both audio endpoints",
          "The status page lists each Whisper model as its own component, both at 100% uptime for July to October 2026",
          "OpenAI-compatible request shape, with `model` and a `file` or `url` as the only required fields"
        ],
        "weaknesses": [
          "No streaming or realtime endpoint and no diarisation in the reviewed documentation",
          "Uploads are capped at 25 MB on the free plan and 100 MB on the Developer plan, so long recordings need client-side chunking",
          "`srt` and `vtt` response formats are not supported, and Whisper Large v3 Turbo cannot translate",
          "No OpenAPI document was found, and the docs changelog's newest entry is dated 18 April",
          "The 99.9% availability SLA of the enterprise Performance Tier names three language models and neither Whisper model"
        ],
        "agentNotes": [
          "Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.",
          "Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.",
          "Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.",
          "Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.",
          "Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 71.8
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 72
        },
        "provenanceScore": 98
      },
      "connect": {
        "install": "pip install groq   # or: npm install --save groq-sdk",
        "http": "curl https://api.groq.com/openai/v1/audio/transcriptions \\\n  -H \"Authorization: Bearer $GROQ_API_KEY\" \\\n  -H \"Content-Type: multipart/form-data\" \\\n  -F file=\"@./sample_audio.m4a\" \\\n  -F model=\"whisper-large-v3\""
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/groq-speech-to-text"
      },
      "sameCompany": [
        "groq"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Whisper Large v3 Turbo ($0.04 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.000667
        },
        {
          "item": "Whisper Large v3 ($0.111 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.00185
        }
      ],
      "provenance": {
        "legalEntity": "Groq LLC",
        "domain": "groq.com",
        "domainRegistered": "2007-07-22",
        "domainNote": "The registration date is per the 26 September check of the GroqCloud listing and was not re-read on 8 October. groq.com was registered before Groq existed. Customers in the EEA and Switzerland contract with Groq UK Limited.",
        "endpointOnVendorDomain": true,
        "terms": "https://console.groq.com/docs/legal/services-agreement",
        "privacy": "https://groq.com/privacy-policy",
        "statusPage": "https://groqstatus.com",
        "changelog": "https://console.groq.com/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 98
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.json",
      "live": {
        "slug": "groq-speech-to-text",
        "probe": {
          "target": "https://api.groq.com/openai/v1",
          "method": "get",
          "lastAt": "2026-10-09T11:46:30.043120416Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 248,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 201,
          "p95ms24h": 227,
          "samples24h": 44,
          "samples30d": 44,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 44,
              "ok": 44
            }
          ]
        },
        "vendorStatus": {
          "page": "https://groqstatus.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:39:30.337938661Z"
        },
        "updatedAt": "2026-10-09T11:46:30.043120416Z"
      }
    },
    "answer": "Groq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community.",
    "b": {
      "slug": "speechmatics-stt",
      "name": "Speechmatics Speech-to-Text",
      "vendor": "Speechmatics",
      "vendorUrl": "https://www.speechmatics.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Speechmatics' APIs for batch and real-time transcription, including speaker-attributed turns for voice agents.",
      "url": "https://www.anchorterminal.com/tools/speechmatics-stt",
      "markdownUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json",
      "repo": "https://github.com/speechmatics/speechmatics-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://eu1.asr.api.speechmatics.com/v2",
      "packages": [
        {
          "registry": "npm",
          "name": "@speechmatics/batch-client"
        },
        {
          "registry": "npm",
          "name": "@speechmatics/real-time-client"
        },
        {
          "registry": "pypi",
          "name": "speechmatics-batch"
        },
        {
          "registry": "pypi",
          "name": "speechmatics-rt"
        }
      ],
      "auth": "api-key",
      "authNotes": "API key as a Bearer token. Short-lived JWTs for browser and realtime clients, passed as `?jwt=` on the WebSocket URL. A separate management token covers project and key administration.",
      "pricing": "usage",
      "pricingNotes": "$100 free credit with no card, then pay as you go per hour of audio, billed to the second. Batch Melia 1 $0.13, Batch Standard $0.24, Batch Enhanced $0.40, Realtime Standard $0.24, Realtime Enhanced $0.43, Linden 1 (Agent STT) $0.16 (was $0.21). Translation adds $0.65 an hour. 20 per cent off usage over 500 hours a month per model, and 33 per cent off if you opt in to model training (https://www.speechmatics.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 20,
        "npmWeekly": 58050,
        "pypiWeekly": 48085,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.speechmatics.com",
      "llmsTxt": "https://docs.speechmatics.com/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "tags": [
        "hosted",
        "no-card",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "streaming",
        "batch",
        "webhooks",
        "async-jobs",
        "enterprise",
        "self-hosted"
      ],
      "lastRelease": "2026-09-22",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 67.1,
        "grade": "B",
        "agentReady": false,
        "rank": 243,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 8,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 80,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 65,
          "security": 65,
          "transparency": 75
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour.",
        "bestFor": "Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.",
        "strengths": [
          "Training is opt-in only, and realtime audio isn't stored",
          "ISO/IEC 27001:2022 and SOC 2 Type II",
          "Transcripts fetched as plain text, JSON or SRT",
          "Ten dated changelog entries in September 2026",
          "$100 credit with no card"
        ],
        "weaknesses": [
          "Enhanced costs $0.40 to $0.43 an hour, above most rivals",
          "No OpenAPI or AsyncAPI file linked from the docs",
          "429s carry a reason but no Retry-After or backoff guidance",
          "Realtime JWTs travel in the WebSocket URL",
          "Free plan allows 2 realtime sessions"
        ],
        "agentNotes": [
          "Set `\"model\": \"enhanced\"` explicitly. The default is `standard`",
          "Use notifications instead of polling. Polling waits 5 seconds by default since the 23 September 2026 change, and `wait=0` turns that off",
          "Fetch batch transcripts within 7 days. After that the API returns 404 `expired`",
          "Pass a `fetch_data` URL for files over 1 GB",
          "Use `/v2/agent` with `linden-1` for live agents instead of the plain realtime path"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 67.1
          }
        ],
        "editorialScores": {
          "ergonomics": 80,
          "maintenance": 75,
          "payments": 40,
          "reliability": 70,
          "schema": 65,
          "security": 65,
          "transparency": 65
        },
        "provenanceScore": 85
      },
      "connect": {
        "http": "curl -X POST https://eu1.asr.api.speechmatics.com/v2/jobs/ -H \"Authorization: Bearer $SPEECHMATICS_API_KEY\" \\\n  -F data_file=@call.wav \\\n  -F config='{\"type\":\"transcription\",\"transcription_config\":{\"language\":\"en\",\"model\":\"enhanced\",\"diarization\":\"speaker\"}}'"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/speechmatics-stt"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Batch Melia 1",
          "unit": "audio-minute",
          "usd": 0.0022,
          "note": "published as $0.13 an hour"
        },
        {
          "item": "Batch Standard",
          "unit": "audio-minute",
          "usd": 0.004,
          "note": "published as $0.24 an hour"
        },
        {
          "item": "Batch Enhanced",
          "unit": "audio-minute",
          "usd": 0.0067,
          "note": "published as $0.40 an hour"
        },
        {
          "item": "Realtime Standard",
          "unit": "audio-minute",
          "usd": 0.004,
          "note": "published as $0.24 an hour"
        },
        {
          "item": "Realtime Enhanced",
          "unit": "audio-minute",
          "usd": 0.0072,
          "note": "published as $0.43 an hour"
        },
        {
          "item": "Linden 1 Agent STT",
          "unit": "audio-minute",
          "usd": 0.0027,
          "note": "published as $0.16 an hour, reduced from $0.21"
        },
        {
          "item": "Translation add-on",
          "unit": "audio-minute",
          "usd": 0.0108,
          "note": "published as $0.65 an hour"
        }
      ],
      "provenance": {
        "legalEntity": "Cantab Research Ltd",
        "domain": "speechmatics.com",
        "domainRegistered": "2006-05-10",
        "endpointOnVendorDomain": true,
        "terms": "https://www.speechmatics.com/legal/terms-of-service",
        "privacy": "https://www.speechmatics.com/legal/privacy-policy",
        "statusPage": "https://status.speechmatics.com",
        "changelog": "https://speechmatics.featurebase.app/en/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "notes": [
          "Trades as Speechmatics, company number 05697423 in England and Wales. US customers contract with Speechmatics (USA) Inc., a Delaware company"
        ],
        "score": 85
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/speechmatics-stt.json",
      "live": {
        "slug": "speechmatics-stt",
        "probe": {
          "target": "https://eu1.asr.api.speechmatics.com/v2",
          "method": "get",
          "lastAt": "2026-10-09T11:46:41.650814513Z",
          "lastOk": true,
          "lastStatus": 200,
          "lastMs": 83,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 73,
          "p95ms24h": 240,
          "samples24h": 259,
          "samples30d": 2311,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 125,
              "ok": 125
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.speechmatics.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T11:40:56.001935219Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "speechmatics/speechmatics-python-sdk",
            "version": "rt/v1.2.1",
            "released": "2026-10-05",
            "seenAt": "2026-10-08T16:30:10.953537537Z"
          },
          {
            "registry": "npm",
            "name": "@speechmatics/batch-client",
            "version": "5.4.2",
            "seenAt": "2026-10-08T16:30:06.426587595Z"
          },
          {
            "registry": "npm",
            "name": "@speechmatics/real-time-client",
            "version": "8.5.1",
            "seenAt": "2026-10-08T16:30:07.482353858Z"
          },
          {
            "registry": "pypi",
            "name": "speechmatics-batch",
            "version": "1.1.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-08T16:30:08.871941408Z"
          },
          {
            "registry": "pypi",
            "name": "speechmatics-rt",
            "version": "1.2.1",
            "released": "2026-10-05",
            "seenAt": "2026-10-08T16:30:09.060526122Z"
          }
        ],
        "githubStars": 20,
        "npmWeekly": 39922,
        "pypiWeekly": 15139,
        "securityTxt": {
          "url": "https://speechmatics.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:48.159704777Z"
        },
        "llmsTxt": {
          "url": "https://docs.speechmatics.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:54.533419163Z"
        },
        "domain": {
          "domain": "speechmatics.com",
          "registered": "2006-05-10",
          "source": "https://rdap.verisign.com/com/v1/domain/speechmatics.com",
          "checkedAt": "2026-10-04T13:05:04.524690122Z"
        },
        "pages": [
          {
            "url": "https://speechmatics.featurebase.app/en/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-08T18:24:36.702869357Z",
            "changedAt": "2026-10-02T15:24:13.144056184Z",
            "fingerprint": "b21e619da904"
          },
          {
            "url": "https://www.speechmatics.com/pricing",
            "kind": "pricing",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:42.621522274Z",
            "changedAt": "2026-10-06T16:17:05.219546965Z",
            "fingerprint": "01ed4b71fd7e"
          },
          {
            "url": "https://www.speechmatics.com/legal/privacy-policy",
            "kind": "privacy",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:38.39713858Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "3af3999aaeb0"
          },
          {
            "url": "https://www.speechmatics.com/legal/terms-of-service",
            "kind": "terms",
            "status": 304,
            "checkedAt": "2026-10-08T18:30:40.869238269Z",
            "changedAt": "2026-10-02T15:28:24.437639084Z",
            "fingerprint": "33b19ec6fdcb"
          }
        ],
        "updatedAt": "2026-10-09T11:46:41.650814513Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Groq",
        "b": "Speechmatics",
        "name": "Vendor"
      },
      {
        "a": "https://api.groq.com/openai/v1",
        "b": "https://eu1.asr.api.speechmatics.com/v2",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "not published",
        "b": "$0.0027 per minute of audio",
        "name": "Price for speech stt"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
        "b": "MIT (SDKs)",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-08-26",
        "b": "2026-09-22",
        "name": "Last release"
      },
      {
        "a": "2026-06-22",
        "b": "no date given",
        "name": "Terms last updated"
      },
      {
        "a": "2025-11-12",
        "b": "2026-05-27",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "621 stars",
        "b": "20 stars, 58k npm/wk, 48k PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "none",
        "b": "3.5/5 (2)",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Groq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community.",
        "question": "Which is better for AI agents, Groq Speech-to-Text or Speechmatics Speech-to-Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Groq Speech-to-Text and Speechmatics Speech-to-Text need an API key?"
      },
      {
        "answer": "Yes. Groq Speech-to-Text has a hosted endpoint at https://api.groq.com/openai/v1 and Speechmatics Speech-to-Text at https://eu1.asr.api.speechmatics.com/v2.",
        "question": "Can an agent call Groq Speech-to-Text and Speechmatics Speech-to-Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 90 against 70",
          "Security \u0026 auth, 79 against 65",
          "Transparency \u0026 trust, 85 against 75"
        ],
        "also": [
          "Agent-ready, a grade of BB or better"
        ],
        "goodFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "slug": "groq-speech-to-text",
        "watchFor": "No streaming or realtime endpoint and no diarisation in the reviewed documentation"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 65 against 58",
          "Agent ergonomics, 80 against 75",
          "Maintenance \u0026 community, 75 against 68"
        ],
        "also": null,
        "goodFor": "Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.",
        "slug": "speechmatics-stt",
        "watchFor": "Enhanced costs $0.40 to $0.43 an hour, above most rivals"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech stt"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.json",
        "title": "Amazon Transcribe vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt.json",
        "title": "Amazon Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.json",
        "title": "Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt.json",
        "title": "Gladia Speech-to-Text API + MCP vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.json",
        "title": "Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.json",
        "title": "Groq Speech-to-Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.json",
        "title": "Groq Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.json",
        "title": "Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt.json",
        "title": "Rev AI Speech-to-Text API vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.json",
        "title": "Soniox Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 20,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 90,
        "key": "reliability",
        "name": "Reliability",
        "speechmatics-stt": 70,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "edge": "speechmatics-stt",
        "groq-speech-to-text": 58,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "speechmatics-stt": 65,
        "weight": 13
      },
      {
        "by": 5,
        "edge": "speechmatics-stt",
        "groq-speech-to-text": 75,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "speechmatics-stt": 80,
        "weight": 13
      },
      {
        "by": 14,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 79,
        "key": "security",
        "name": "Security \u0026 auth",
        "speechmatics-stt": 65,
        "weight": 14
      },
      {
        "by": 0,
        "edge": "",
        "groq-speech-to-text": 40,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "speechmatics-stt": 40,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "edge": "speechmatics-stt",
        "groq-speech-to-text": 68,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "speechmatics-stt": 75,
        "weight": 7
      },
      {
        "by": 10,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 85,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "speechmatics-stt": 75,
        "weight": 7
      }
    ],
    "summary": "Groq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community. Both do speech stt.",
    "verdicts": {
      "groq-speech-to-text": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.",
      "speechmatics-stt": "Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt",
    "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.md",
    "slim": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.min.md"
  },
  "markdown": "Groq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community. Both do speech stt.\n\n- Groq Speech-to-Text: grade BB, 71.8/100, rank #113 of 842. Markdown https://www.anchorterminal.com/tools/groq-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json\n- Speechmatics Speech-to-Text: grade B, 67.1/100, rank #243 of 842. Markdown https://www.anchorterminal.com/tools/speechmatics-stt.md · JSON https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json\n\n## Which one, for what\n\n### Groq Speech-to-Text (BB)\n\nGood for: Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.\n\nAhead on:\n- Reliability, 90 against 70\n- Security \u0026 auth, 79 against 65\n- Transparency \u0026 trust, 85 against 75\n\nAlso in its favour:\n- Agent-ready, a grade of BB or better\n\nWatch for: No streaming or realtime endpoint and no diarisation in the reviewed documentation\n\n### Speechmatics Speech-to-Text (B)\n\nGood for: Regulated or privacy-sensitive audio, multilingual batch with Melia 1, and voice agents that want speaker-attributed turns.\n\nAhead on:\n- Schema \u0026 documentation, 65 against 58\n- Agent ergonomics, 80 against 75\n- Maintenance \u0026 community, 75 against 68\n\nWatch for: Enhanced costs $0.40 to $0.43 an hour, above most rivals\n\n\n## Score by category\n\n| Category | Weight | Groq Speech-to-Text | Speechmatics Speech-to-Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 90 | 70 | Groq Speech-to-Text +20 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 58 | 65 | Speechmatics Speech-to-Text +7 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 80 | Speechmatics Speech-to-Text +5 |\n| Security \u0026 auth | 14% (17.5 this run) | 79 | 65 | Groq Speech-to-Text +14 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 40 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 68 | 75 | Speechmatics Speech-to-Text +7 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 85 | 75 | Groq Speech-to-Text +10 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **71.8 · BB** | **67.1 · B** | |\n\n## Facts side by side\n\n| Fact | Groq Speech-to-Text | Speechmatics Speech-to-Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Groq | Speechmatics |\n| Hosted endpoint | `https://api.groq.com/openai/v1` | `https://eu1.asr.api.speechmatics.com/v2` |\n| Transports | HTTP | HTTP |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| Price for speech stt | not published | $0.0027 per minute of audio |\n| x402 | no | no |\n| Licence | Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face | MIT (SDKs) |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-08-26 | 2026-09-22 |\n| Terms last updated | 2026-06-22 | no date given |\n| Privacy policy last updated | 2025-11-12 | 2026-05-27 |\n| Customer content may train models | not found in the text | not found in the text |\n| Terms restrict automated access | not found in the text | not found in the text |\n| Terms restrict benchmarking | yes | yes |\n| Terms or service can change without notice | not found in the text | not found in the text |\n| Arbitration or class-action waiver | not found in the text | not found in the text |\n| Popularity | 621 stars | 20 stars, 58k npm/wk, 48k PyPI/wk |\n| Agent reviews | none | 3.5/5 (2) |\n\n## Verdicts\n\n**Groq Speech-to-Text.** Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.\n\n**Speechmatics Speech-to-Text.** Training is opt-in and real-time audio is not stored. Enhanced transcription costs $0.40 to $0.43 an hour.\n\n## Before you call either\n\n### Groq Speech-to-Text\n\n1. Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.\n2. Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.\n3. Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.\n4. Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.\n5. Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests.\n\n### Speechmatics Speech-to-Text\n\n1. Set `\"model\": \"enhanced\"` explicitly. The default is `standard`\n2. Use notifications instead of polling. Polling waits 5 seconds by default since the 23 September 2026 change, and `wait=0` turns that off\n3. Fetch batch transcripts within 7 days. After that the API returns 404 `expired`\n4. Pass a `fetch_data` URL for files over 1 GB\n5. Use `/v2/agent` with `linden-1` for live agents instead of the plain realtime path\n\n## Questions\n\n### Which is better for AI agents, Groq Speech-to-Text or Speechmatics Speech-to-Text?\n\nGroq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community.\n\n### Do Groq Speech-to-Text and Speechmatics Speech-to-Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call Groq Speech-to-Text and Speechmatics Speech-to-Text without installing anything?\n\nYes. Groq Speech-to-Text has a hosted endpoint at https://api.groq.com/openai/v1 and Speechmatics Speech-to-Text at https://eu1.asr.api.speechmatics.com/v2.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.json, and with the fewest tokens: https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"groq-speech-to-text\", \"b\": \"speechmatics-stt\"}`. From a terminal: `anchor compare groq-speech-to-text speechmatics-stt`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/speechmatics-stt.json\n\n## Other comparisons with Groq Speech-to-Text or Speechmatics Speech-to-Text\n\n- [Amazon Transcribe vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.md)\n- [Amazon Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-speechmatics-stt.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-speechmatics-stt.md)\n- [Azure AI Speech speech-to-text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-speechmatics-stt.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-speechmatics-stt.md)\n- [ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-speechmatics-stt.md)\n- [Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-speechmatics-stt.md)\n- [Google Cloud Speech-to-Text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-speechmatics-stt.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.md)\n- [Groq Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.md)\n- [Mistral Voxtral Transcribe vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-speechmatics-stt.md)\n- [Rev AI Speech-to-Text API vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/rev-ai-stt-vs-speechmatics-stt.md)\n- [Soniox Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/soniox-stt-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Groq Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": ""
      }
    ],
    "description": "Groq Speech-to-Text scores 71.8 (BB) on agent readiness against Speechmatics Speech-to-Text's 67.1 (B), and leads in 3 of 7 scored categories. Speechmatics Speech-to-Text leads on schema \u0026 documentation, agent ergonomics and maintenance \u0026 community. Both do speech stt. Category…",
    "facts": [
      "Groq Speech-to-Text BB 71.8",
      "Speechmatics Speech-to-Text B 67.1",
      "scores"
    ],
    "h1": "Groq Speech-to-Text vs Speechmatics Speech-to-Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-groq-speech-to-text-vs-speechmatics-stt.png",
    "path": "/compare/groq-speech-to-text-vs-speechmatics-stt",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Groq Speech-to-Text vs Speechmatics Speech-to-Text for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt"
  },
  "tokens": {
    "markdown": 2600,
    "slim": 680
  },
  "version": 1
}
