{
  "data": {
    "a": {
      "slug": "deepgram-voice-agent",
      "name": "Deepgram Voice Agent API",
      "vendor": "Deepgram",
      "vendorUrl": "https://deepgram.com",
      "kind": "http-api",
      "category": "voice-agents",
      "summary": "One WebSocket that runs Deepgram STT (Flux or Nova-3), a managed or bring-your-own LLM and Deepgram or third-party TTS, with turn-taking, barge-in and function calling.",
      "url": "https://www.anchorterminal.com/tools/deepgram-voice-agent",
      "markdownUrl": "https://www.anchorterminal.com/tools/deepgram-voice-agent.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/deepgram-voice-agent.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/deepgram-voice-agent.json",
      "repo": "https://github.com/deepgram/deepgram-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "streamable-http",
        "stdio",
        "sse"
      ],
      "remoteUrl": "https://agent.deepgram.com/v1/agent",
      "packages": [
        {
          "registry": "npm",
          "name": "@deepgram/sdk"
        },
        {
          "registry": "pypi",
          "name": "deepgram-sdk"
        },
        {
          "registry": "pypi",
          "name": "deepctl"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Token \u003ckey\u003e` header on REST and WebSocket calls. Short-lived JWTs (30-second TTL) from the token endpoint for browsers. The `dg` CLI MCP server uses `dg login` credentials or `DEEPGRAM_API_KEY`. The docs MCP needs no key. BYO LLM and TTS providers take their own keys in the `endpoint.headers` of the Settings message.",
      "pricing": "usage",
      "pricingNotes": "Billed per minute of WebSocket connection time. Standard $0.075 a minute pay as you go ($0.068 Growth), Standard with your own TTS $0.065, your own LLM and TTS $0.050 ($0.041 Growth). Advanced tier (larger LLMs such as GPT-5 or Claude Sonnet) $0.163, or $0.122 with your own TTS. $200 free credit with no card. Carrier costs are extra (https://deepgram.com/pricing).",
      "priceSummary": "Pay per use",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing. Card or prepaid credits only (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 468,
        "npmWeekly": 1123798,
        "pypiWeekly": 805026,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://developers.deepgram.com/docs/voice-agent",
      "llmsTxt": "https://developers.deepgram.com/llms.txt",
      "capabilities": [
        "voice.agent",
        "voice.pipeline",
        "voice.tools"
      ],
      "tags": [
        "hosted",
        "no-card",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "mcp",
        "streaming",
        "pipeline",
        "enterprise",
        "self-hosted"
      ],
      "lastRelease": "2026-09-30",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 68.1,
        "grade": "B",
        "agentReady": false,
        "rank": 213,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 3,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 67,
          "maintenance": 85,
          "payments": 40,
          "reliability": 55,
          "schema": 90,
          "security": 71,
          "transparency": 77
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "$0.075 a minute all in on Standard, $0.050 with your own LLM and TTS, $200 of credit with no card. No phone numbers or SIP, so you bridge Twilio or another carrier yourself.",
        "bestFor": "Developers who already run telephony (Twilio, Amazon Connect, Genesys, AudioCodes) and want a cheap managed pipeline with their choice of LLM.",
        "strengths": [
          "$0.075 a minute all in on Standard, $0.050 with your own LLM and TTS, $200 of credit with no card",
          "AsyncAPI description of the agent socket plus a public OpenAPI file",
          "Role-based API keys with expiry and 30-second browser tokens",
          "SDKs in Python, JavaScript, Java, Go, C# and React, updated in September 2026",
          "Function-call hold that waits for a confirmed user turn before irreversible tools run"
        ],
        "weaknesses": [
          "No phone numbers or SIP, so you bridge Twilio or another carrier yourself",
          "Fifteen status incidents since July, four of them an hour or longer on parts the agent uses",
          "Call audio is kept for model improvement unless `mip_opt_out` is set",
          "No SLA terms on self-serve plans",
          "Sessions end at 2 hours"
        ],
        "agentNotes": [
          "Set `mip_opt_out` in the Settings message to keep call audio out of training",
          "Use the function-call hold for anything irreversible, since replies can start before the user's turn is confirmed",
          "Back off exponentially on 429, a pay-as-you-go project gets 45 concurrent sockets",
          "Pin the LLM version, an unpinned Gemini route failed for 1.5 hours on 21 July 2026",
          "Mint browser tokens with `ttl_seconds`, not `ttl`, or they expire after 30 seconds"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 3,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "B",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 68.1
          }
        ],
        "editorialScores": {
          "ergonomics": 67,
          "maintenance": 85,
          "payments": 40,
          "reliability": 55,
          "schema": 90,
          "security": 71,
          "transparency": 70
        },
        "provenanceScore": 84
      },
      "connect": {
        "http": "curl https://agent.deepgram.com/v1/agent/settings/think/models \\\n  -H \"Authorization: Token $DEEPGRAM_API_KEY\"",
        "claudeCode": "claude mcp add deepgram-docs --transport http https://api.dx.deepgram.com/kapa/mcp",
        "config": {
          "mcpServers": {
            "deepgram": {
              "args": [
                "mcp"
              ],
              "command": "dg",
              "env": {
                "DEEPGRAM_API_KEY": "${DEEPGRAM_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/voice.agent",
        "tool": "https://letme.dev/deepgram-voice-agent"
      },
      "sameCompany": [
        "deepgram-stt",
        "deepgram-tts"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Standard",
          "unit": "call-minute",
          "usd": 0.075,
          "note": "pay as you go, Deepgram STT, managed LLM and TTS included, carrier extra"
        },
        {
          "item": "Standard with own TTS",
          "unit": "call-minute",
          "usd": 0.065
        },
        {
          "item": "Own LLM and TTS",
          "unit": "call-minute",
          "usd": 0.05,
          "note": "STT and orchestration only"
        },
        {
          "item": "Advanced",
          "unit": "call-minute",
          "usd": 0.163,
          "note": "larger LLMs, all-in except carrier"
        },
        {
          "item": "Advanced with own TTS",
          "unit": "call-minute",
          "usd": 0.122
        }
      ],
      "provenance": {
        "legalEntity": "Deepgram, Inc.",
        "domain": "deepgram.com",
        "domainRegistered": "2016-01-28",
        "endpointOnVendorDomain": true,
        "terms": "https://deepgram.com/terms",
        "privacy": "https://deepgram.com/privacy",
        "statusPage": "https://status.deepgram.com",
        "changelog": "https://developers.deepgram.com/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 84
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/deepgram-voice-agent.json",
      "live": {
        "slug": "deepgram-voice-agent",
        "probe": {
          "target": "https://agent.deepgram.com/v1/agent",
          "method": "get",
          "lastAt": "2026-10-09T10:42:41.164937584Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 460,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 456,
          "p95ms24h": 582,
          "samples24h": 260,
          "samples30d": 2300,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-05",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-06",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-07",
              "probes": 272,
              "ok": 272
            },
            {
              "date": "2026-10-08",
              "probes": 268,
              "ok": 268
            },
            {
              "date": "2026-10-09",
              "probes": 114,
              "ok": 114
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.deepgram.com",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-09T10:41:33.17390694Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "deepgram/deepgram-python-sdk",
            "version": "v7.12.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-08T16:08:28.56403361Z"
          },
          {
            "registry": "npm",
            "name": "@deepgram/sdk",
            "version": "5.14.0",
            "seenAt": "2026-10-08T16:08:23.193501884Z"
          },
          {
            "registry": "pypi",
            "name": "deepctl",
            "version": "0.3.2",
            "released": "2026-10-05",
            "seenAt": "2026-10-08T16:08:26.708417462Z"
          },
          {
            "registry": "pypi",
            "name": "deepgram-sdk",
            "version": "7.12.0",
            "released": "2026-10-02",
            "seenAt": "2026-10-08T16:08:26.555624685Z"
          }
        ],
        "githubStars": 469,
        "npmWeekly": 1133092,
        "pypiWeekly": 784290,
        "securityTxt": {
          "url": "https://deepgram.com/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-08T15:38:51.488361424Z"
        },
        "llmsTxt": {
          "url": "https://developers.deepgram.com/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-08T14:00:22.707684961Z"
        },
        "domain": {
          "domain": "deepgram.com",
          "registered": "2016-01-28",
          "source": "https://rdap.verisign.com/com/v1/domain/deepgram.com",
          "checkedAt": "2026-10-04T13:03:28.939824686Z"
        },
        "updatedAt": "2026-10-09T10:42:41.164937584Z"
      }
    },
    "answer": "Deepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories.",
    "b": {
      "slug": "gemini-live",
      "name": "Gemini Live API",
      "vendor": "Google",
      "vendorUrl": "https://ai.google.dev",
      "kind": "http-api",
      "category": "voice-agents",
      "summary": "Google's Live API runs real-time spoken conversations with Gemini audio-to-audio models over a stateful WebSocket. It takes audio, images and text, speaks back, and supports interruptions, function calling and Google Search grounding.",
      "url": "https://www.anchorterminal.com/tools/gemini-live",
      "markdownUrl": "https://www.anchorterminal.com/tools/gemini-live.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/gemini-live.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/gemini-live.json",
      "repo": "https://github.com/google-gemini/gemini-live-api-examples",
      "license": "Proprietary service under the Gemini API Additional Terms of Service. The Python and JavaScript SDKs are Apache-2.0",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://generativelanguage.googleapis.com/v1beta",
      "packages": [
        {
          "registry": "pypi",
          "name": "google-genai"
        },
        {
          "registry": "npm",
          "name": "@google/genai"
        }
      ],
      "auth": "api-key",
      "authNotes": "Self-serve API key from Google AI Studio, tied to a Google Cloud project. Keys created since 28 May 2026 are authorisation keys bound to a service account and restricted to the Gemini API, and unrestricted standard keys are rejected. The raw WebSocket guide passes the key as a `key` query parameter. For browsers and phones, a backend mints an ephemeral token (preview) at `POST /v1beta/auth_tokens`, single use by default, 1 minute to start a session and 30 minutes to use it, optionally locked to a model and configuration.",
      "pricing": "freemium",
      "pricingNotes": "Free tier on the Live models with no billing account, where content is used to improve Google products. Paid tier per 1M tokens on `gemini-3.8-live`, the extended thinking model and `gemini-3.1-flash-live-preview`. Text in $0.75, audio in $3.00 (about $0.005 a minute), image or video in $1.00, text out $4.50, audio out $12.00 (about $0.018 a minute). Google Search grounding $14 per 1,000 queries after 5,000 free a month. Each turn re-bills the whole session context. Paid tier needs a linked billing account and a $5 minimum prepayment (https://ai.google.dev/gemini-api/docs/pricing, checked 2026-10-08).",
      "priceSummary": "$14 / 1k req",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402, MPP or L402 in the Live API docs, the pricing page or the billing guide (checked 2026-10-08).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": null,
        "npmWeekly": 29486957,
        "pypiWeekly": 34128301,
        "asOf": "2026-10-08"
      },
      "docsUrl": "https://ai.google.dev/gemini-api/docs/live-api",
      "llmsTxt": "https://ai.google.dev/gemini-api/docs/llms.txt",
      "capabilities": [
        "voice.agent",
        "voice.speech-to-speech",
        "voice.tools"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "speech-to-speech",
        "websocket",
        "streaming",
        "python",
        "typescript",
        "llms-txt",
        "api-key",
        "function-calling"
      ],
      "lastRelease": "2026-09-15",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 57.1,
        "grade": "C",
        "agentReady": false,
        "rank": 560,
        "ranked": true,
        "rankOf": 842,
        "categoryRank": 7,
        "methodology": "0.4",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 69,
          "maintenance": 83,
          "payments": 40,
          "reliability": 41,
          "schema": 66,
          "security": 48,
          "transparency": 72
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-08"
        },
        "negative": 0,
        "verdict": "A speech-to-speech API with per-token prices published, a free tier and a generally available model, `gemini-3.8-live`, since 15 September 2026. The documented raw WebSocket connection carries the API key in the URL, connections reset about every 10 minutes, and no SLA or readable incident history was found for the Developer API.",
        "bestFor": "Teams building their own voice or vision assistant on a single speech-to-speech model with function calling and Search grounding, who can run a backend for tokens and audio transport.",
        "strengths": [
          "`gemini-3.8-live` went generally available on 15 September 2026 with asynchronous function calling as the default and three response scheduling modes.",
          "Ephemeral tokens can be single use, expire in 30 minutes by default and be locked to a model and session configuration.",
          "Prices are public per million tokens with per-minute equivalents, $0.005 a minute of audio in and $0.018 a minute of audio out.",
          "Session resumption handles, a GoAway message with `timeLeft` and sliding-window context compression are documented for long sessions.",
          "Free tier on the Live models with no billing account, and paid-tier content isn't used to improve Google products per the pricing page."
        ],
        "weaknesses": [
          "The WebSocket guide authenticates with the API key as a `key` query parameter, and ephemeral tokens as an `access_token` query parameter.",
          "The status page at aistudio.google.com/status renders in the browser only, so no incident history could be read, and no SLA was found for the Developer API.",
          "No machine-readable contract for the WebSocket messages was found. The Discovery document types only the setup and token schemas.",
          "The capabilities, best practices and API reference pages still say the Live API is in preview, and the docs give both 70 and 99 supported languages.",
          "No telephony. Phone calls need a partner such as Voximplant, LiveKit or Pipecat, and concurrent session limits are shown only inside AI Studio."
        ],
        "agentNotes": [
          "Enable `sessionResumption` and keep the newest handle. Connections end after about 10 minutes, and handles stay valid for 2 hours.",
          "Set `contextWindowCompression` with a sliding window. Without it audio sessions stop at 15 minutes and audio with video at 2 minutes, and every turn re-bills the whole context.",
          "On `gemini-3.8-live` function calls are non-blocking by default. Set `behavior: BLOCKING` if the model must wait for the tool response.",
          "Send 16 kHz 16-bit PCM in 20 to 40 ms chunks and discard buffered playback when `interrupted` is true.",
          "Keep the API key on a server and send it through the SDK. Give browsers an ephemeral token from `POST /v1beta/auth_tokens`, locked with `liveConnectConstraints`."
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 0,
        "avgRating": 0,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.4",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 57.1
          }
        ],
        "editorialScores": {
          "ergonomics": 69,
          "maintenance": 83,
          "payments": 40,
          "reliability": 41,
          "schema": 66,
          "security": 48,
          "transparency": 55
        },
        "provenanceScore": 89
      },
      "connect": {
        "install": "pip install google-genai   # or: npm i @google/genai",
        "http": "curl -X POST \"https://generativelanguage.googleapis.com/v1beta/auth_tokens\" \\\n  -H \"x-goog-api-key: $GEMINI_API_KEY\" -H \"Content-Type: application/json\" \\\n  -d '{\"uses\": 1, \"liveConnectConstraints\": {\"model\": \"models/gemini-3.8-live\", \"config\": {\"sessionResumption\": {}, \"responseModalities\": [\"AUDIO\"]}}}'"
      },
      "letme": {
        "capability": "https://letme.dev/voice.agent",
        "tool": "https://letme.dev/gemini-live"
      },
      "sameCompany": [
        "gemini-api",
        "gemini-embedding",
        "vertex-ai-tuning",
        "google-model-armor",
        "google-imagen",
        "google-veo",
        "google-lyria",
        "google-speech-to-text",
        "google-adk",
        "google-secret-manager",
        "google-weather-api",
        "chrome-devtools-mcp",
        "google-maps-platform",
        "google-cloud-translation",
        "google-calendar-api",
        "firebase-cloud-messaging",
        "google-drive-api",
        "gemini-cli",
        "google-search-console",
        "google-ads-api",
        "google-forms",
        "google-sheets-api",
        "gmail-api"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "gemini-3.8-live audio input",
          "unit": "audio-minute",
          "usd": 0.005,
          "note": "$3.00 per 1M audio tokens, charged while the API is listening"
        },
        {
          "item": "gemini-3.8-live audio output",
          "unit": "audio-minute",
          "usd": 0.018,
          "note": "$12.00 per 1M audio tokens"
        },
        {
          "item": "gemini-3.8-live text input",
          "unit": "1m-tokens",
          "usd": 0.75
        },
        {
          "item": "gemini-3.8-live text output",
          "unit": "1m-tokens",
          "usd": 4.5,
          "note": "includes thinking tokens and transcription text"
        },
        {
          "item": "Google Search grounding",
          "unit": "1k-requests",
          "usd": 14,
          "note": "after 5,000 free search queries a month shared across Gemini 3.x models"
        }
      ],
      "provenance": {
        "legalEntity": "Google LLC",
        "domain": "google.com",
        "domainRegistered": "1997-09-15",
        "endpointOnVendorDomain": true,
        "terms": "https://ai.google.dev/gemini-api/terms",
        "privacy": "https://policies.google.com/privacy",
        "statusPage": "https://aistudio.google.com/status",
        "changelog": "https://ai.google.dev/gemini-api/docs/changelog",
        "securityTxt": "valid",
        "checked": "2026-10-08",
        "notes": [
          "The WebSocket and the token endpoint are on generativelanguage.googleapis.com, Google's API domain. The docs are on ai.google.dev.",
          "RDAP for google.com gives a registration date of 1997-09-15.",
          "The Gemini API Additional Terms of Service are effective 23 March 2026 and incorporate the Google APIs Terms of Service. Paid services fall under Google's Data Processing Addendum for products where Google is a data processor.",
          "The Google Privacy Policy read on 8 October 2026 is effective 1 October 2026. It is Google's general policy, and the terms page is where API data use is set out.",
          "www.google.com/.well-known/security.txt expires 2030-04-01 and points to g.co/vulnz and the vulnerability reward programme. ai.google.dev has no security.txt of its own (404).",
          "aistudio.google.com/status answered 200 with a page drawn by script, so no component or incident history was read."
        ],
        "score": 89
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/gemini-live.json",
      "live": {
        "slug": "gemini-live",
        "probe": {
          "target": "https://generativelanguage.googleapis.com/v1beta",
          "method": "get",
          "lastAt": "2026-10-09T10:42:44.067853492Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 17,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 25,
          "p95ms24h": 49,
          "samples24h": 33,
          "samples30d": 33,
          "days": [
            {
              "date": "2026-10-09",
              "probes": 33,
              "ok": 33
            }
          ]
        },
        "updatedAt": "2026-10-09T10:42:44.067853492Z"
      }
    },
    "facts": [
      {
        "a": "HTTP API",
        "b": "HTTP API",
        "name": "Kind"
      },
      {
        "a": "Deepgram",
        "b": "Google",
        "name": "Vendor"
      },
      {
        "a": "https://agent.deepgram.com/v1/agent",
        "b": "https://generativelanguage.googleapis.com/v1beta",
        "name": "Hosted endpoint"
      },
      {
        "a": "HTTP, Streamable HTTP, stdio, SSE (legacy)",
        "b": "HTTP",
        "name": "Transports"
      },
      {
        "a": "API key",
        "b": "API key",
        "name": "Auth"
      },
      {
        "a": "Pay per use",
        "b": "Freemium",
        "name": "Pricing"
      },
      {
        "a": "no",
        "b": "no",
        "name": "x402"
      },
      {
        "a": "MIT (SDKs)",
        "b": "Proprietary service under the Gemini API Additional Terms of Service. The Python and JavaScript SDKs are Apache-2.0",
        "name": "Licence"
      },
      {
        "a": "no",
        "b": "no",
        "name": "Read-only variant documented"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "llms.txt"
      },
      {
        "a": "2026-09-30",
        "b": "2026-09-15",
        "name": "Last release"
      },
      {
        "a": "2026-08-06",
        "b": "2026-04-28",
        "name": "Terms last updated"
      },
      {
        "a": "2021-10-26",
        "b": "2026-10-01",
        "name": "Privacy policy last updated"
      },
      {
        "a": "yes, with an opt-out",
        "b": "yes",
        "name": "Customer content may train models"
      },
      {
        "a": "not found in the text",
        "b": "yes",
        "name": "Terms restrict automated access"
      },
      {
        "a": "yes",
        "b": "yes",
        "name": "Terms restrict benchmarking"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Terms or service can change without notice"
      },
      {
        "a": "yes",
        "b": "not found in the text",
        "name": "Arbitration or class-action waiver"
      },
      {
        "a": "468 stars, 1.1M npm/wk, 805k PyPI/wk",
        "b": "29.5M npm/wk, 34.1M PyPI/wk",
        "name": "Popularity"
      },
      {
        "a": "3/5 (2)",
        "b": "none",
        "name": "Agent reviews"
      }
    ],
    "faq": [
      {
        "answer": "Deepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories.",
        "question": "Which is better for AI agents, Deepgram Voice Agent API or Gemini Live API?"
      },
      {
        "answer": "Both need an API key.",
        "question": "Do Deepgram Voice Agent API and Gemini Live API need an API key?"
      },
      {
        "answer": "Yes. Deepgram Voice Agent API has a hosted endpoint at https://agent.deepgram.com/v1/agent and Gemini Live API at https://generativelanguage.googleapis.com/v1beta.",
        "question": "Can an agent call Deepgram Voice Agent API and Gemini Live API without installing anything?"
      }
    ],
    "goodFor": [
      {
        "aheadOn": [
          "Reliability, 55 against 41",
          "Schema \u0026 documentation, 90 against 66",
          "Security \u0026 auth, 71 against 48",
          "Transparency \u0026 trust, 77 against 72"
        ],
        "also": [
          "Runs on your own machine",
          "Free to start without a card"
        ],
        "goodFor": "Developers who already run telephony (Twilio, Amazon Connect, Genesys, AudioCodes) and want a cheap managed pipeline with their choice of LLM.",
        "slug": "deepgram-voice-agent",
        "watchFor": "No phone numbers or SIP, so you bridge Twilio or another carrier yourself"
      },
      {
        "aheadOn": null,
        "also": null,
        "goodFor": "Teams building their own voice or vision assistant on a single speech-to-speech model with function calling and Search grounding, who can run a backend for tokens and audio transport.",
        "slug": "gemini-live",
        "watchFor": "The WebSocket guide authenticates with the API key as a `key` query parameter, and ephemeral tokens as an `access_token` query parameter."
      }
    ],
    "job": {
      "capability": "voice.agent",
      "name": "Voice agent"
    },
    "others": [
      {
        "json": "https://www.anchorterminal.com/compare/bland-ai-vs-deepgram-voice-agent.json",
        "title": "Bland AI API + MCP vs Deepgram Voice Agent API",
        "url": "https://www.anchorterminal.com/compare/bland-ai-vs-deepgram-voice-agent"
      },
      {
        "json": "https://www.anchorterminal.com/compare/bland-ai-vs-gemini-live.json",
        "title": "Bland AI API + MCP vs Gemini Live API",
        "url": "https://www.anchorterminal.com/compare/bland-ai-vs-gemini-live"
      },
      {
        "json": "https://www.anchorterminal.com/compare/bolna-vs-deepgram-voice-agent.json",
        "title": "Bolna API + MCP vs Deepgram Voice Agent API",
        "url": "https://www.anchorterminal.com/compare/bolna-vs-deepgram-voice-agent"
      },
      {
        "json": "https://www.anchorterminal.com/compare/bolna-vs-gemini-live.json",
        "title": "Bolna API + MCP vs Gemini Live API",
        "url": "https://www.anchorterminal.com/compare/bolna-vs-gemini-live"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-elevenlabs-agents.json",
        "title": "Deepgram Voice Agent API vs ElevenLabs Agents API + MCP",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-elevenlabs-agents"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-hume-evi.json",
        "title": "Deepgram Voice Agent API vs Hume EVI (Empathic Voice Interface)",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-hume-evi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-retell-ai.json",
        "title": "Deepgram Voice Agent API vs Retell AI API + MCP",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-retell-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-synthflow.json",
        "title": "Deepgram Voice Agent API vs Synthflow API + MCP",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-synthflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-ultravox.json",
        "title": "Deepgram Voice Agent API vs Ultravox Realtime API",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-ultravox"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vapi.json",
        "title": "Deepgram Voice Agent API vs Vapi API + MCP",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vapi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vocily.json",
        "title": "Deepgram Voice Agent API vs Vocily AI",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vocily"
      },
      {
        "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vogent.json",
        "title": "Deepgram Voice Agent API vs Vogent API",
        "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vogent"
      },
      {
        "json": "https://www.anchorterminal.com/compare/elevenlabs-agents-vs-gemini-live.json",
        "title": "ElevenLabs Agents API + MCP vs Gemini Live API",
        "url": "https://www.anchorterminal.com/compare/elevenlabs-agents-vs-gemini-live"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-hume-evi.json",
        "title": "Gemini Live API vs Hume EVI (Empathic Voice Interface)",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-hume-evi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-retell-ai.json",
        "title": "Gemini Live API vs Retell AI API + MCP",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-retell-ai"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-synthflow.json",
        "title": "Gemini Live API vs Synthflow API + MCP",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-synthflow"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-ultravox.json",
        "title": "Gemini Live API vs Ultravox Realtime API",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-ultravox"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-vapi.json",
        "title": "Gemini Live API vs Vapi API + MCP",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-vapi"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-vocily.json",
        "title": "Gemini Live API vs Vocily AI",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-vocily"
      },
      {
        "json": "https://www.anchorterminal.com/compare/gemini-live-vs-vogent.json",
        "title": "Gemini Live API vs Vogent API",
        "url": "https://www.anchorterminal.com/compare/gemini-live-vs-vogent"
      }
    ],
    "scores": [
      {
        "by": 14,
        "deepgram-voice-agent": 55,
        "edge": "deepgram-voice-agent",
        "gemini-live": 41,
        "key": "reliability",
        "name": "Reliability",
        "weight": 16
      },
      {
        "key": "performance",
        "name": "Performance",
        "pending": true,
        "weight": 10
      },
      {
        "by": 24,
        "deepgram-voice-agent": 90,
        "edge": "deepgram-voice-agent",
        "gemini-live": 66,
        "key": "schema",
        "name": "Schema \u0026 documentation",
        "weight": 13
      },
      {
        "by": 2,
        "deepgram-voice-agent": 67,
        "edge": "gemini-live",
        "gemini-live": 69,
        "key": "ergonomics",
        "name": "Agent ergonomics",
        "weight": 13
      },
      {
        "by": 23,
        "deepgram-voice-agent": 71,
        "edge": "deepgram-voice-agent",
        "gemini-live": 48,
        "key": "security",
        "name": "Security \u0026 auth",
        "weight": 14
      },
      {
        "by": 0,
        "deepgram-voice-agent": 40,
        "edge": "",
        "gemini-live": 40,
        "key": "payments",
        "name": "Payments \u0026 pricing",
        "weight": 10
      },
      {
        "key": "tasks",
        "name": "Task success",
        "pending": true,
        "weight": 10
      },
      {
        "by": 2,
        "deepgram-voice-agent": 85,
        "edge": "deepgram-voice-agent",
        "gemini-live": 83,
        "key": "maintenance",
        "name": "Maintenance \u0026 community",
        "weight": 7
      },
      {
        "by": 5,
        "deepgram-voice-agent": 77,
        "edge": "deepgram-voice-agent",
        "gemini-live": 72,
        "key": "transparency",
        "name": "Transparency \u0026 trust",
        "weight": 7
      }
    ],
    "summary": "Deepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories. Both do voice agent.",
    "verdicts": {
      "deepgram-voice-agent": "$0.075 a minute all in on Standard, $0.050 with your own LLM and TTS, $200 of credit with no card. No phone numbers or SIP, so you bridge Twilio or another carrier yourself.",
      "gemini-live": "A speech-to-speech API with per-token prices published, a free tier and a generally available model, `gemini-3.8-live`, since 15 September 2026. The documented raw WebSocket connection carries the API key in the URL, connections reset about every 10 minutes, and no SLA or readable incident history was found for the Developer API."
    }
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live",
    "json": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live.md",
    "slim": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live.min.md"
  },
  "markdown": "Deepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories. Both do voice agent.\n\n- Deepgram Voice Agent API: grade B, 68.1/100, rank #213 of 842. Markdown https://www.anchorterminal.com/tools/deepgram-voice-agent.md · JSON https://www.anchorterminal.com/api/v1/tools/deepgram-voice-agent.json\n- Gemini Live API: grade C, 57.1/100, rank #560 of 842. Markdown https://www.anchorterminal.com/tools/gemini-live.md · JSON https://www.anchorterminal.com/api/v1/tools/gemini-live.json\n\n## Which one, for what\n\n### Deepgram Voice Agent API (B)\n\nGood for: Developers who already run telephony (Twilio, Amazon Connect, Genesys, AudioCodes) and want a cheap managed pipeline with their choice of LLM.\n\nAhead on:\n- Reliability, 55 against 41\n- Schema \u0026 documentation, 90 against 66\n- Security \u0026 auth, 71 against 48\n- Transparency \u0026 trust, 77 against 72\n\nAlso in its favour:\n- Runs on your own machine\n- Free to start without a card\n\nWatch for: No phone numbers or SIP, so you bridge Twilio or another carrier yourself\n\n### Gemini Live API (C)\n\nGood for: Teams building their own voice or vision assistant on a single speech-to-speech model with function calling and Search grounding, who can run a backend for tokens and audio transport.\n\nWatch for: The WebSocket guide authenticates with the API key as a `key` query parameter, and ephemeral tokens as an `access_token` query parameter.\n\n\n## Score by category\n\n| Category | Weight | Deepgram Voice Agent API | Gemini Live API | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 55 | 41 | Deepgram Voice Agent API +14 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 90 | 66 | Deepgram Voice Agent API +24 |\n| Agent ergonomics | 13% (16.2 this run) | 67 | 69 | Gemini Live API +2 |\n| Security \u0026 auth | 14% (17.5 this run) | 71 | 48 | Deepgram Voice Agent API +23 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 40 | 40 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 85 | 83 | Deepgram Voice Agent API +2 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 77 | 72 | Deepgram Voice Agent API +5 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **68.1 · B** | **57.1 · C** | |\n\n## Facts side by side\n\n| Fact | Deepgram Voice Agent API | Gemini Live API |\n| --- | --- | --- |\n| Kind | HTTP API | HTTP API |\n| Vendor | Deepgram | Google |\n| Hosted endpoint | `https://agent.deepgram.com/v1/agent` | `https://generativelanguage.googleapis.com/v1beta` |\n| Transports | HTTP, Streamable HTTP, stdio, SSE (legacy) | HTTP |\n| Auth | API key | API key |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | MIT (SDKs) | Proprietary service under the Gemini API Additional Terms of Service. The Python and JavaScript SDKs are Apache-2.0 |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| Last release | 2026-09-30 | 2026-09-15 |\n| Terms last updated | 2026-08-06 | 2026-04-28 |\n| Privacy policy last updated | 2021-10-26 | 2026-10-01 |\n| Customer content may train models | yes, with an opt-out | yes |\n| Terms restrict automated access | not found in the text | yes |\n| Terms restrict benchmarking | yes | yes |\n| Terms or service can change without notice | yes | not found in the text |\n| Arbitration or class-action waiver | yes | not found in the text |\n| Popularity | 468 stars, 1.1M npm/wk, 805k PyPI/wk | 29.5M npm/wk, 34.1M PyPI/wk |\n| Agent reviews | 3/5 (2) | none |\n\n## Verdicts\n\n**Deepgram Voice Agent API.** $0.075 a minute all in on Standard, $0.050 with your own LLM and TTS, $200 of credit with no card. No phone numbers or SIP, so you bridge Twilio or another carrier yourself.\n\n**Gemini Live API.** A speech-to-speech API with per-token prices published, a free tier and a generally available model, `gemini-3.8-live`, since 15 September 2026. The documented raw WebSocket connection carries the API key in the URL, connections reset about every 10 minutes, and no SLA or readable incident history was found for the Developer API.\n\n## Before you call either\n\n### Deepgram Voice Agent API\n\n1. Set `mip_opt_out` in the Settings message to keep call audio out of training\n2. Use the function-call hold for anything irreversible, since replies can start before the user's turn is confirmed\n3. Back off exponentially on 429, a pay-as-you-go project gets 45 concurrent sockets\n4. Pin the LLM version, an unpinned Gemini route failed for 1.5 hours on 21 July 2026\n5. Mint browser tokens with `ttl_seconds`, not `ttl`, or they expire after 30 seconds\n\n### Gemini Live API\n\n1. Enable `sessionResumption` and keep the newest handle. Connections end after about 10 minutes, and handles stay valid for 2 hours.\n2. Set `contextWindowCompression` with a sliding window. Without it audio sessions stop at 15 minutes and audio with video at 2 minutes, and every turn re-bills the whole context.\n3. On `gemini-3.8-live` function calls are non-blocking by default. Set `behavior: BLOCKING` if the model must wait for the tool response.\n4. Send 16 kHz 16-bit PCM in 20 to 40 ms chunks and discard buffered playback when `interrupted` is true.\n5. Keep the API key on a server and send it through the SDK. Give browsers an ephemeral token from `POST /v1beta/auth_tokens`, locked with `liveConnectConstraints`.\n\n## Questions\n\n### Which is better for AI agents, Deepgram Voice Agent API or Gemini Live API?\n\nDeepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories.\n\n### Do Deepgram Voice Agent API and Gemini Live API need an API key?\n\nBoth need an API key.\n\n### Can an agent call Deepgram Voice Agent API and Gemini Live API without installing anything?\n\nYes. Deepgram Voice Agent API has a hosted endpoint at https://agent.deepgram.com/v1/agent and Gemini Live API at https://generativelanguage.googleapis.com/v1beta.\n\n\n## For agents\n\n- This comparison as JSON: https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live.json, and with the fewest tokens: https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live.min.md\n- Over MCP at https://www.anchorterminal.com/mcp (no key): `compare_tools {\"a\": \"deepgram-voice-agent\", \"b\": \"gemini-live\"}`. From a terminal: `anchor compare deepgram-voice-agent gemini-live`\n- Each listing in full: https://www.anchorterminal.com/api/v1/tools/deepgram-voice-agent.json and https://www.anchorterminal.com/api/v1/tools/gemini-live.json\n\n## Other comparisons with Deepgram Voice Agent API or Gemini Live API\n\n- [Bland AI API + MCP vs Deepgram Voice Agent API](https://www.anchorterminal.com/compare/bland-ai-vs-deepgram-voice-agent.md)\n- [Bland AI API + MCP vs Gemini Live API](https://www.anchorterminal.com/compare/bland-ai-vs-gemini-live.md)\n- [Bolna API + MCP vs Deepgram Voice Agent API](https://www.anchorterminal.com/compare/bolna-vs-deepgram-voice-agent.md)\n- [Bolna API + MCP vs Gemini Live API](https://www.anchorterminal.com/compare/bolna-vs-gemini-live.md)\n- [Deepgram Voice Agent API vs ElevenLabs Agents API + MCP](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-elevenlabs-agents.md)\n- [Deepgram Voice Agent API vs Hume EVI (Empathic Voice Interface)](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-hume-evi.md)\n- [Deepgram Voice Agent API vs Retell AI API + MCP](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-retell-ai.md)\n- [Deepgram Voice Agent API vs Synthflow API + MCP](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-synthflow.md)\n- [Deepgram Voice Agent API vs Ultravox Realtime API](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-ultravox.md)\n- [Deepgram Voice Agent API vs Vapi API + MCP](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vapi.md)\n- [Deepgram Voice Agent API vs Vocily AI](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vocily.md)\n- [Deepgram Voice Agent API vs Vogent API](https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-vogent.md)\n- [ElevenLabs Agents API + MCP vs Gemini Live API](https://www.anchorterminal.com/compare/elevenlabs-agents-vs-gemini-live.md)\n- [Gemini Live API vs Hume EVI (Empathic Voice Interface)](https://www.anchorterminal.com/compare/gemini-live-vs-hume-evi.md)\n- [Gemini Live API vs Retell AI API + MCP](https://www.anchorterminal.com/compare/gemini-live-vs-retell-ai.md)\n- [Gemini Live API vs Synthflow API + MCP](https://www.anchorterminal.com/compare/gemini-live-vs-synthflow.md)\n- [Gemini Live API vs Ultravox Realtime API](https://www.anchorterminal.com/compare/gemini-live-vs-ultravox.md)\n- [Gemini Live API vs Vapi API + MCP](https://www.anchorterminal.com/compare/gemini-live-vs-vapi.md)\n- [Gemini Live API vs Vocily AI](https://www.anchorterminal.com/compare/gemini-live-vs-vocily.md)\n- [Gemini Live API vs Vogent API](https://www.anchorterminal.com/compare/gemini-live-vs-vogent.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-09",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Deepgram Voice Agent API vs Gemini Live API",
        "url": ""
      }
    ],
    "description": "Deepgram Voice Agent API scores 68.1 (B) on agent readiness against Gemini Live API's 57.1 (C), and leads in 5 of 7 scored categories. Both do voice agent. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Deepgram Voice Agent API B 68.1",
      "Gemini Live API C 57.1",
      "scores"
    ],
    "h1": "Deepgram Voice Agent API vs Gemini Live API",
    "image": "https://www.anchorterminal.com/assets/og/compare-deepgram-voice-agent-vs-gemini-live.png",
    "path": "/compare/deepgram-voice-agent-vs-gemini-live",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Deepgram Voice Agent API vs Gemini Live API for AI agents",
    "toc": null,
    "updated": "2026-10-09",
    "url": "https://www.anchorterminal.com/compare/deepgram-voice-agent-vs-gemini-live"
  },
  "tokens": {
    "markdown": 2450,
    "slim": 730
  },
  "version": 1
}
