{
  "data": {
    "a": {
      "slug": "groq-speech-to-text",
      "name": "Groq Speech-to-Text",
      "vendor": "Groq",
      "vendorUrl": "https://groq.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "Groq's hosted speech-to-text API. It runs OpenAI's Whisper Large v3 and Whisper Large v3 Turbo on OpenAI-compatible transcription and translation endpoints, for uploaded files or audio URLs, with a half-price batch mode.",
      "url": "https://www.anchorterminal.com/tools/groq-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json",
      "repo": "https://github.com/groq/groq-python",
      "license": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.groq.com/openai/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "groq"
        },
        {
          "registry": "npm",
          "name": "groq-sdk"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve bearer key from the GroqCloud console, created inside a project. Projects carry their own rate limits per model, usage data and request logs (https://console.groq.com/docs/projects).",
      "pricing": "freemium",
      "pricingNotes": "$0.04 per audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with a 10-second minimum a request and 50% off through the Batch API (https://console.groq.com/docs/models). The free plan needs no card, so an agent's owner can start without a contract. The Developer plan is postpaid by card, US bank account or SEPA debit.",
      "priceSummary": "Freemium",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the speech-to-text guide, the API reference or the billing pages (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 621,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://console.groq.com/docs/speech-to-text",
      "llmsTxt": "https://console.groq.com/llms.txt",
      "capabilities": [
        "speech.stt",
        "speech.batch",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "fast",
        "free-tier",
        "no-card",
        "llms-txt",
        "python",
        "typescript",
        "batch",
        "openai-compatible"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 71.8,
        "grade": "BB",
        "agentReady": true,
        "rank": 117,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 4,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 85
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.",
        "bestFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "strengths": [
          "Published prices of $0.04 an audio hour for Whisper Large v3 Turbo and $0.111 for Whisper Large v3, with 50% off through the Batch API",
          "Free plan with no card at 20 requests a minute, 2,000 a day and 7,200 audio seconds an hour on both models",
          "Inputs and outputs are not retained by default, and zero data retention is a console setting that covers both audio endpoints",
          "The status page lists each Whisper model as its own component, both at 100% uptime for July to October 2026",
          "OpenAI-compatible request shape, with `model` and a `file` or `url` as the only required fields"
        ],
        "weaknesses": [
          "No streaming or realtime endpoint and no diarisation in the reviewed documentation",
          "Uploads are capped at 25 MB on the free plan and 100 MB on the Developer plan, so long recordings need client-side chunking",
          "`srt` and `vtt` response formats are not supported, and Whisper Large v3 Turbo cannot translate",
          "No OpenAPI document was found, and the docs changelog's newest entry is dated 18 April",
          "The 99.9% availability SLA of the enterprise Performance Tier names three language models and neither Whisper model"
        ],
        "agentNotes": [
          "Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.",
          "Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.",
          "Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.",
          "Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.",
          "Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 71.8
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 68,
          "payments": 40,
          "reliability": 90,
          "schema": 58,
          "security": 79,
          "transparency": 72
        },
        "provenanceScore": 98
      },
      "connect": {
        "install": "pip install groq   # or: npm install --save groq-sdk",
        "http": "curl https://api.groq.com/openai/v1/audio/transcriptions \\\n  -H \"Authorization: Bearer $GROQ_API_KEY\" \\\n  -H \"Content-Type: multipart/form-data\" \\\n  -F file=\"@./sample_audio.m4a\" \\\n  -F model=\"whisper-large-v3\""
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/groq-speech-to-text"
      },
      "sameCompany": [
        "groq"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Whisper Large v3 Turbo ($0.04 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.000667
        },
        {
          "item": "Whisper Large v3 ($0.111 an audio hour)",
          "unit": "audio-minute",
          "usd": 0.00185
        }
      ],
      "provenance": {
        "legalEntity": "Groq LLC",
        "domain": "groq.com",
        "domainRegistered": "2007-07-22",
        "domainNote": "The registration date is per the 26 September check of the GroqCloud listing and was not re-read on 8 October. groq.com was registered before Groq existed. Customers in the EEA and Switzerland contract with Groq UK Limited.",
        "endpointOnVendorDomain": true,
        "terms": "https://console.groq.com/docs/legal/services-agreement",
        "privacy": "https://groq.com/privacy-policy",
        "statusPage": "https://groqstatus.com",
        "changelog": "https://console.groq.com/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "score": 98
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/groq-speech-to-text.json",
      "live": {
        "slug": "groq-speech-to-text",
        "probe": {
          "target": "https://api.groq.com/openai/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:07:07.598634731Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 177,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 199,
          "p95ms24h": 284,
          "samples24h": 201,
          "samples30d": 201,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 169,
              "ok": 169
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://groqstatus.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-10T03:02:13.557858059Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "groq/groq-python",
            "version": "v1.7.0",
            "released": "2026-08-26",
            "seenAt": "2026-10-09T16:56:57.125960169Z"
          },
          {
            "registry": "npm",
            "name": "groq-sdk",
            "version": "1.6.0",
            "seenAt": "2026-10-09T16:56:56.300671531Z"
          },
          {
            "registry": "pypi",
            "name": "groq",
            "version": "1.7.0",
            "released": "2026-08-26",
            "seenAt": "2026-10-09T16:56:56.108246835Z"
          }
        ],
        "githubStars": 621,
        "npmWeekly": 884716,
        "pypiWeekly": 4040988,
        "securityTxt": {
          "url": "https://groq.com/.well-known/security.txt",
          "state": "valid",
          "checkedAt": "2026-10-09T15:40:10.425828811Z"
        },
        "llmsTxt": {
          "url": "https://console.groq.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-09T14:02:04.377441459Z"
        },
        "pages": [
          {
            "url": "https://console.groq.com/docs/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-09T18:34:17.894319038Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "ac2cb86ef05a"
          },
          {
            "url": "https://groq.com/privacy-policy",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-09T18:39:59.465247347Z",
            "changedAt": "2026-10-09T18:39:59.465247347Z",
            "fingerprint": "1ae299b9f26b"
          },
          {
            "url": "https://console.groq.com/docs/legal/services-agreement",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-09T18:34:22.225273307Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "a0b93cf713d4"
          }
        ],
        "updatedAt": "2026-10-10T03:07:07.598634731Z"
      }
    },
    "answer": "OpenAI Speech to Text and Groq Speech-to-Text score within a point of each other on agent readiness, 72.4 (BB) and 71.8 (BB). Groq Speech-to-Text leads on reliability, payments \u0026 pricing and transparency \u0026 trust.",
    "b": {
      "slug": "openai-speech-to-text",
      "name": "OpenAI Speech to Text",
      "vendor": "OpenAI",
      "vendorUrl": "https://openai.com",
      "kind": "model",
      "category": "speech-to-text",
      "summary": "OpenAI's speech-to-text API. It transcribes uploaded audio files through `/v1/audio/transcriptions`, translates recordings into English through `/v1/audio/translations`, and transcribes live audio in Realtime transcription sessions over WebSocket or WebRTC.",
      "url": "https://www.anchorterminal.com/tools/openai-speech-to-text",
      "markdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json",
      "repo": "https://github.com/openai/openai-python",
      "license": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
      "transports": [
        "http",
        "websocket"
      ],
      "remoteUrl": "https://api.openai.com/v1",
      "packages": [
        {
          "registry": "pypi",
          "name": "openai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Bearer API key created by a person in the platform console (https://platform.openai.com/settings/organization/api-keys). Projects can carry a model allowlist or denylist and an IP allowlist, and Admin API keys are a separate credential that cannot call the audio endpoints (https://developers.openai.com/api/docs/guides/admin-apis).",
      "pricing": "usage",
      "pricingNotes": "$0.0045 an audio minute for `gpt-transcribe` and $0.017 for `gpt-live-transcribe`, billed from prepaid credits (https://developers.openai.com/api/docs/pricing). The rate limits guide names a Free tier with a $100 monthly usage limit, and the `gpt-transcribe` model page lists limits only from the Build tier, which needs $5 of credit purchases. Whether a new account can transcribe without paying was not established.",
      "priceSummary": "Pay per use",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the transcription guides, the endpoint reference, the pricing page or the OpenAPI document (checked 2026-10-09).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 31785,
        "npmWeekly": null,
        "pypiWeekly": null,
        "asOf": "2026-10-09"
      },
      "docsUrl": "https://developers.openai.com/api/docs/guides/speech-to-text",
      "llmsTxt": "https://developers.openai.com/llms.txt",
      "openapi": "https://github.com/openai/openai-openapi/blob/main/openapi.yaml",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "tags": [
        "official",
        "hosted",
        "model",
        "streaming",
        "diarisation",
        "llms-txt",
        "openapi",
        "python",
        "typescript",
        "go",
        "java"
      ],
      "lastRelease": "2026-08-26",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 72.4,
        "grade": "BB",
        "agentReady": true,
        "rank": 106,
        "ranked": true,
        "rankOf": 950,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-09"
        },
        "negative": 0,
        "verdict": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.",
        "bestFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "strengths": [
          "`gpt-transcribe` is priced at $0.0045 an audio minute on a public page, with per-tier request limits of 5,000, 10,000 and 30,000 a minute",
          "The data controls page lists `/v1/audio/transcriptions` and `/v1/audio/translations` with no training, no abuse-monitoring retention and no stored application state",
          "A public OpenAPI 3.1 document, llms.txt and a Markdown twin of every docs page cover the audio endpoints",
          "The status page has an Audio component, shown at 100% uptime for July to October 2026",
          "Guide examples cover JavaScript, Python, Go, Java, C#, Ruby, a CLI and curl, and `file` and `model` are the only required fields"
        ],
        "weaknesses": [
          "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027",
          "The two named replacements return no speaker labels, word timestamps, `srt` or `vtt` output or English translation in the reviewed documentation",
          "Uploads stop at 25 MB, the caller splits longer recordings, and the `gpt-transcribe` model page marks the Batch API as not supported",
          "The Markdown twin of the endpoint reference lists the response fields and omits the request parameters, and its first example names a deprecated model",
          "openai.com answered our reader with a bot check, so the service terms, privacy policy, sub-processor list and any SLA were not read"
        ],
        "agentNotes": [
          "Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.",
          "Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.",
          "For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.",
          "Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.",
          "On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "BB",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 72.4
          }
        ],
        "editorialScores": {
          "ergonomics": 78,
          "maintenance": 75,
          "payments": 20,
          "reliability": 80,
          "schema": 88,
          "security": 86,
          "transparency": 62
        },
        "provenanceScore": 59
      },
      "connect": {
        "install": "pip install openai",
        "http": "curl --request POST \\\n  --url https://api.openai.com/v1/audio/transcriptions \\\n  --header \"Authorization: Bearer $OPENAI_API_KEY\" \\\n  --header 'Content-Type: multipart/form-data' \\\n  --form file=@/path/to/file/audio.mp3 \\\n  --form model=gpt-transcribe"
      },
      "letme": {
        "capability": "https://letme.dev/speech.stt",
        "tool": "https://letme.dev/openai-speech-to-text"
      },
      "sameCompany": [
        "openai-api",
        "openai-embeddings",
        "openai-guardrails",
        "openai-moderation",
        "openai-image-api",
        "openai-sora",
        "openai-realtime",
        "openai-agents-sdk",
        "openai-decisions-api",
        "openai-codex"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gpt-transcribe",
          "unit": "audio-minute",
          "usd": 0.0045
        },
        {
          "item": "gpt-live-transcribe (live audio)",
          "unit": "audio-minute",
          "usd": 0.017
        },
        {
          "item": "whisper-1 (deprecated)",
          "unit": "audio-minute",
          "usd": 0.006
        },
        {
          "item": "gpt-4o-transcribe-diarize (deprecated, estimated from token prices)",
          "unit": "audio-minute",
          "usd": 0.006
        }
      ],
      "provenance": {
        "legalEntity": "",
        "domain": "openai.com",
        "domainRegistered": "",
        "domainNote": "openai.com answered our researcher with a bot check on 9 October 2026, so the terms and privacy policy were not read on that day. The links are the two documents OpenAI's other listings here carry. openai.com answers our policy reader with HTTP 403 as well, so neither document has been read and both are recorded as unreadable. security.txt is PGP-signed with Bugcrowd and email contacts and has no Expires field.",
        "endpointOnVendorDomain": true,
        "terms": "https://openai.com/policies/services-agreement/",
        "privacy": "https://openai.com/policies/privacy-policy/",
        "statusPage": "https://status.openai.com",
        "changelog": "https://developers.openai.com/api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-09",
        "score": 59
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/openai-speech-to-text.json",
      "live": {
        "slug": "openai-speech-to-text",
        "probe": {
          "target": "https://api.openai.com/v1",
          "method": "get",
          "lastAt": "2026-10-10T03:07:15.369334284Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 133,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 137,
          "p95ms24h": 172,
          "samples24h": 117,
          "samples30d": 117,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 85,
              "ok": 85
            },
            {
              "date": "2026-10-10",
              "probes": 32,
              "ok": 32
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.openai.com",
          "indicator": "minor",
          "summary": "Partial System Degradation",
          "checkedAt": "2026-10-10T03:02:29.613651511Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "openai/openai-python",
            "version": "v3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.877366615Z"
          },
          {
            "registry": "pypi",
            "name": "openai",
            "version": "3.27.0",
            "released": "2026-10-09",
            "seenAt": "2026-10-09T17:10:31.73110612Z"
          }
        ],
        "githubStars": 31787,
        "pypiWeekly": 74761714,
        "updatedAt": "2026-10-10T03:07:15.369334284Z"
      }
    },
    "facts": [
      {
        "a": "Model API",
        "b": "Model API",
        "name": "Kind"
      },
      {
        "a": "Groq",
        "b": "OpenAI",
        "name": "Vendor"
      },
      {
        "a": "https://api.groq.com/openai/v1",
        "b": "https://api.openai.com/v1",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP",
        "b": "HTTP, websocket",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Freemium",
        "b": "Pay per use",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face",
        "b": "Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-08-26",
        "b": "2026-08-26",
        "name": "Last release"
      },
      {
        "a": "2026-06-22",
        "b": "couldn't be read",
        "name": "Terms last updated"
      },
      {
        "a": "2025-11-12",
        "b": "couldn't be read",
        "name": "Privacy policy last updated"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "couldn't be read",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "not found in the text",
        "b": "couldn't be read",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "621 stars",
        "b": "32k stars",
        "name": "Popularity"
      }
    ],
    "faq": [
      {
        "answer": "OpenAI Speech to Text and Groq Speech-to-Text score within a point of each other on agent readiness, 72.4 (BB) and 71.8 (BB). Groq Speech-to-Text leads on reliability, payments \u0026 pricing and transparency \u0026 trust.",
        "question": "Which is better for AI agents, Groq Speech-to-Text or OpenAI Speech to Text?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Groq Speech-to-Text and OpenAI Speech to Text need an API key?"
      },
      {
        "answer": "Yes. Groq Speech-to-Text has a hosted endpoint at https://api.groq.com/openai/v1 and OpenAI Speech to Text at https://api.openai.com/v1.",
        "question": "Can an agent call Groq Speech-to-Text and OpenAI Speech to Text without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 90 against 80",
          "Payments \u0026 pricing, 40 against 20",
          "Transparency \u0026 trust, 85 against 61"
        ],
        "also": [
          "Free to start without a card"
        ],
        "goodFor": "Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.",
        "slug": "groq-speech-to-text",
        "watchFor": "No streaming or realtime endpoint and no diarisation in the reviewed documentation"
      },
      {
        "aheadOn": [
          "Schema \u0026 documentation, 88 against 58",
          "Security \u0026 auth, 86 against 79",
          "Maintenance \u0026 community, 75 against 68"
        ],
        "also": null,
        "goodFor": "Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.",
        "slug": "openai-speech-to-text",
        "watchFor": "`whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027"
      }
    ],
    "job": {
      "capability": "speech.stt",
      "name": "Speech-to-text"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.json",
        "title": "Amazon Transcribe vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.json",
        "title": "Amazon Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.json",
        "title": "AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Azure AI Speech speech-to-text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text.json",
        "title": "Cartesia Ink vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.json",
        "title": "Cartesia Ink vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.json",
        "title": "Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.json",
        "title": "ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.json",
        "title": "Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs Groq Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.json",
        "title": "Google Cloud Speech-to-Text vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.json",
        "title": "Groq Speech-to-Text vs Mistral Voxtral Transcribe",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.json",
        "title": "Groq Speech-to-Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.json",
        "title": "Groq Speech-to-Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.json",
        "title": "Groq Speech-to-Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.json",
        "title": "Mistral Voxtral Transcribe vs OpenAI Speech to Text",
        "url": "https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.json",
        "title": "OpenAI Speech to Text vs Rev AI Speech-to-Text API",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.json",
        "title": "OpenAI Speech to Text vs Soniox Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt"
      },
      {
        "json": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.json",
        "title": "OpenAI Speech to Text vs Speechmatics Speech-to-Text",
        "url": "https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt"
      }
    ],
    "scores": [
      {
        "by": 10,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 90,
        "key": "reliability",
        "name": "Reliability",
        "openai-speech-to-text": 80,
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 30,
        "edge": "openai-speech-to-text",
        "groq-speech-to-text": 58,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "openai-speech-to-text": 88,
        "weight": 13
      },
      {
        "by": 3,
        "edge": "openai-speech-to-text",
        "groq-speech-to-text": 75,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "openai-speech-to-text": 78,
        "weight": 13
      },
      {
        "by": 7,
        "edge": "openai-speech-to-text",
        "groq-speech-to-text": 79,
        "key": "security",
        "name": "Security \u0026 auth",
        "openai-speech-to-text": 86,
        "weight": 14
      },
      {
        "by": 20,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 40,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "openai-speech-to-text": 20,
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 7,
        "edge": "openai-speech-to-text",
        "groq-speech-to-text": 68,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "openai-speech-to-text": 75,
        "weight": 7
      },
      {
        "by": 24,
        "edge": "groq-speech-to-text",
        "groq-speech-to-text": 85,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "openai-speech-to-text": 61,
        "weight": 7
      }
    ],
    "summary": "OpenAI Speech to Text and Groq Speech-to-Text score within a point of each other on agent readiness, 72.4 (BB) and 71.8 (BB). Groq Speech-to-Text leads on reliability, payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.",
    "verdicts": {
      "groq-speech-to-text": "Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.",
      "openai-speech-to-text": "`gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text",
    "json": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.md",
    "slim": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.min.md"
  },
  "markdown": "OpenAI Speech to Text and Groq Speech-to-Text score within a point of each other on agent readiness, 72.4 (BB) and 71.8 (BB). Groq Speech-to-Text leads on reliability, payments \u0026 pricing and transparency \u0026 trust. Both do speech-to-text.\n\n- Groq Speech-to-Text: grade BB, 71.8/100, rank #117 of 950. Markdown https://www.anchorterminal.com/tools/groq-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json\n- OpenAI Speech to Text: grade BB, 72.4/100, rank #106 of 950. Markdown https://www.anchorterminal.com/tools/openai-speech-to-text.md · JSON https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n- Best speech-to-text APIs for AI agents: https://www.anchorterminal.com/best/speech-to-text/index.md\n- All 91 stt comparisons: https://www.anchorterminal.com/compare/speech-to-text/index.md\n\n## Which one, for what\n\n### Groq Speech-to-Text (BB)\n\nGood for: Suited to cheap, fast transcription of recorded files in many languages, and to agents that already hold a Groq key or use an OpenAI-compatible client.\n\nAhead on:\n- Reliability, 90 against 80\n- Payments \u0026 pricing, 40 against 20\n- Transparency \u0026 trust, 85 against 61\n\nAlso in its favour:\n- Free to start without a card\n\nWatch for: No streaming or realtime endpoint and no diarisation in the reviewed documentation\n\n### OpenAI Speech to Text (BB)\n\nGood for: Suited to plain transcription of recorded files at a low price a minute and to teams already holding an OpenAI key.\n\nAhead on:\n- Schema \u0026 documentation, 88 against 58\n- Security \u0026 auth, 86 against 79\n- Maintenance \u0026 community, 75 against 68\n\nWatch for: `whisper-1`, `gpt-4o-transcribe`, `gpt-4o-mini-transcribe` and `gpt-4o-transcribe-diarize` were deprecated on 26 August 2026 and shut down on 26 February 2027\n\n\n## Score by category\n\n| Category | Weight | Groq Speech-to-Text | OpenAI Speech to Text | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 90 | 80 | Groq Speech-to-Text +10 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 58 | 88 | OpenAI Speech to Text +30 |\n| Agent ergonomics | 13% (16.2 this run) | 75 | 78 | OpenAI Speech to Text +3 |\n| Security \u0026 auth | 14% (17.5 this run) | 79 | 86 | OpenAI Speech to Text +7 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 20 | Groq Speech-to-Text +20 |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 68 | 75 | OpenAI Speech to Text +7 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 85 | 61 | Groq Speech-to-Text +24 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **71.8 · BB** | **72.4 · BB** | |\n\n## Facts side by side\n\n| Fact | Groq Speech-to-Text | OpenAI Speech to Text |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Groq | OpenAI |\n| Hosted endpoint | `https://api.groq.com/openai/v1` | `https://api.openai.com/v1` |\n| Transports | HTTP | HTTP, websocket |\n| Auth | API key | API key |\n| Pricing | Freemium | Pay per use |\n| x402 | no | no |\n| Licence | Proprietary hosted service under the Groq Services Agreement. The SDKs are Apache-2.0 and the Whisper weights are published by OpenAI on Hugging Face | Proprietary hosted service. The service terms were not read (see open questions). The Python SDK is Apache-2.0 and the OpenAPI document is MIT |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-08-26 | 2026-08-26 |\n| Terms last updated | 2026-06-22 | couldn't be read |\n| Privacy policy last updated | 2025-11-12 | couldn't be read |\n| Customer content may train models | not found in the text | couldn't be read |\n| Terms restrict automated access | not found in the text | couldn't be read |\n| Terms restrict benchmarking | yes | couldn't be read |\n| Terms or service can change without notice | not found in the text | couldn't be read |\n| Arbitration or class-action waiver | not found in the text | couldn't be read |\n| Popularity | 621 stars | 32k stars |\n\n## Verdicts\n\n**Groq Speech-to-Text.** Whisper Large v3 Turbo costs $0.04 an audio hour and Whisper Large v3 $0.111, with a no-card free plan and zero data retention as a self-serve setting. There is no streaming endpoint, no diarisation and no subtitle output, and uploads stop at 25 MB on the free plan and 100 MB on the Developer plan.\n\n**OpenAI Speech to Text.** `gpt-transcribe` costs $0.0045 an audio minute, and the audio endpoints keep no abuse-monitoring logs or application state. Speaker labels, timestamps, subtitles and translation exist only on `whisper-1` and `gpt-4o-transcribe-diarize`, which shut down on 26 February 2027 with no named replacement for those functions.\n\n## Before you call either\n\n### Groq Speech-to-Text\n\n1. Send `whisper-large-v3-turbo` for transcription and `whisper-large-v3` for translation to English. The translations endpoint does not accept Turbo.\n2. Pass `url` instead of `file` for audio over 25 MB, and split anything over the plan's size limit into overlapping chunks before sending.\n3. Set `response_format` to `verbose_json` before asking for `timestamp_granularities[]`. Word timestamps add latency, segment timestamps do not.\n4. Every request is billed as at least 10 seconds of audio, so join very short clips where the task allows.\n5. Read `retry-after` on a 429 and back off. Audio limits count seconds an hour and a day as well as requests.\n\n### OpenAI Speech to Text\n\n1. Send `gpt-transcribe` to `POST /v1/audio/transcriptions` for recorded files. Use `languages` (a list), not `language`, and never send both.\n2. Keep each upload at 25 MB or less. Split longer audio between sentences and pass the previous chunk's text in `prompt`.\n3. For speaker labels send `gpt-4o-transcribe-diarize` with `response_format=diarized_json` and `chunking_strategy=auto` for audio over 30 seconds. Plan for its shutdown on 26 February 2027.\n4. Word timestamps, `srt`, `vtt` and `/v1/audio/translations` need `whisper-1`, which cannot stream and shuts down on the same date.\n5. On 429 or 503 wait at least `Retry-After` when present, then back off with jitter. Do not retry `credit_balance_exhausted` or spend-limit errors.\n\n## Questions\n\n### Which is better for AI agents, Groq Speech-to-Text or OpenAI Speech to Text?\n\nOpenAI Speech to Text and Groq Speech-to-Text score within a point of each other on agent readiness, 72.4 (BB) and 71.8 (BB). Groq Speech-to-Text leads on reliability, payments \u0026 pricing and transparency \u0026 trust.\n\n### Do Groq Speech-to-Text and OpenAI Speech to Text need an API key?\n\nBoth need an API key.\n\n### Can an agent call Groq Speech-to-Text and OpenAI Speech to Text without installing anything?\n\nYes. Groq Speech-to-Text has a hosted endpoint at https://api.groq.com/openai/v1 and OpenAI Speech to Text at https://api.openai.com/v1.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.json, and with the fewest tokens: https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"groq-speech-to-text\", \"b\": \"openai-speech-to-text\"}`. From a terminal: `anchor compare groq-speech-to-text openai-speech-to-text`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/groq-speech-to-text.json and https://www.anchorterminal.com/api/v1/tools/openai-speech-to-text.json\n\n## Other comparisons with Groq Speech-to-Text or OpenAI Speech to Text\n\n- [Amazon Transcribe vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-groq-speech-to-text.md)\n- [Amazon Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/amazon-transcribe-vs-openai-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-groq-speech-to-text.md)\n- [AssemblyAI Speech-to-Text (Universal) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/assemblyai-stt-vs-openai-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-groq-speech-to-text.md)\n- [Azure AI Speech speech-to-text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/azure-speech-to-text-vs-openai-speech-to-text.md)\n- [Cartesia Ink vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-groq-speech-to-text.md)\n- [Cartesia Ink vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/cartesia-ink-stt-vs-openai-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-groq-speech-to-text.md)\n- [Deepgram Speech-to-Text (Nova-3, Flux) vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/deepgram-stt-vs-openai-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-groq-speech-to-text.md)\n- [ElevenLabs Scribe Speech to Text API vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/elevenlabs-scribe-vs-openai-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/gladia-stt-vs-groq-speech-to-text.md)\n- [Gladia Speech-to-Text API + MCP vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/gladia-stt-vs-openai-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs Groq Speech-to-Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-groq-speech-to-text.md)\n- [Google Cloud Speech-to-Text vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/google-speech-to-text-vs-openai-speech-to-text.md)\n- [Groq Speech-to-Text vs Mistral Voxtral Transcribe](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-mistral-voxtral-transcribe.md)\n- [Groq Speech-to-Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-rev-ai-stt.md)\n- [Groq Speech-to-Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-soniox-stt.md)\n- [Groq Speech-to-Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/groq-speech-to-text-vs-speechmatics-stt.md)\n- [Mistral Voxtral Transcribe vs OpenAI Speech to Text](https://www.anchorterminal.com/compare/mistral-voxtral-transcribe-vs-openai-speech-to-text.md)\n- [OpenAI Speech to Text vs Rev AI Speech-to-Text API](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-rev-ai-stt.md)\n- [OpenAI Speech to Text vs Soniox Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-soniox-stt.md)\n- [OpenAI Speech to Text vs Speechmatics Speech-to-Text](https://www.anchorterminal.com/compare/openai-speech-to-text-vs-speechmatics-stt.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Groq Speech-to-Text vs OpenAI Speech to Text",
        "url": ""
      }
    ],
    "description": "OpenAI Speech to Text and Groq Speech-to-Text score within a point of each other for speech-to-text, 72.4 and 71.8 out of 100. Prices, MCP, x402, uptime and agent notes side by side.",
    "facts": [
      "Groq Speech-to-Text BB 71.8",
      "OpenAI Speech to Text BB 72.4",
      "scores"
    ],
    "h1": "Groq Speech-to-Text vs OpenAI Speech to Text",
    "image": "https://www.anchorterminal.com/assets/og/compare-groq-speech-to-text-vs-openai-speech-to-text.png",
    "path": "/compare/groq-speech-to-text-vs-openai-speech-to-text",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Groq Speech-to-Text vs OpenAI Speech to Text for AI agents (2026)",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/groq-speech-to-text-vs-openai-speech-to-text"
  },
  "tokens": {
    "markdown": 2900,
    "slim": 780
  },
  "version": 1
}
