{
  "data": {
    "a": {
      "slug": "fish-audio-voice-cloning",
      "name": "Fish Audio Voice Cloning API",
      "vendor": "Fish Audio",
      "vendorUrl": "https://fish.audio",
      "kind": "model",
      "category": "voice-cloning",
      "summary": "Fish Audio's API creates reusable voice clones from audio samples and generates candidate voices from text prompts.",
      "url": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning",
      "markdownUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/fish-audio-voice-cloning.json",
      "repo": "https://github.com/fishaudio/fish-audio-python",
      "license": "Apache-2.0 (Python SDK)",
      "transports": [
        "http"
      ],
      "remoteUrl": "https://api.fish.audio",
      "packages": [
        {
          "registry": "pypi",
          "name": "fish-audio-sdk"
        },
        {
          "registry": "npm",
          "name": "fish-audio"
        }
      ],
      "auth": "api-key",
      "authNotes": "`Authorization: Bearer` API key. Voice design also needs a `model: voice-design-1` header.",
      "pricing": "usage",
      "pricingNotes": "Prepaid pay as you go with no subscription. No published fee for creating a voice model. Speech from a clone costs $15 per million UTF-8 bytes on `s2.1-pro`, or $0 on `s2.1-pro-free` under fair use. Voice design is $0.01 per successful request whatever the number of candidates. Concurrency rises with total prepaid spend. Separate app plans (Plus $11, Pro $75 a month) set voice slots in the web app (https://docs.fish.audio/developer-guide/models-pricing/pricing-and-rate-limits).",
      "priceSummary": "$0.01 / call",
      "where": "hosted",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": null,
      "popularity": {
        "githubStars": 218,
        "npmWeekly": 6392,
        "pypiWeekly": 37344,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://docs.fish.audio/developer-guide/core-features/creating-models",
      "llmsTxt": "https://docs.fish.audio/llms.txt",
      "openapi": "https://docs.fish.audio/api-reference/openapi.json",
      "capabilities": [
        "voice.clone",
        "voice.design",
        "speech.tts"
      ],
      "tags": [
        "hosted",
        "usage-priced",
        "prepaid",
        "python",
        "typescript",
        "openapi",
        "llms-txt",
        "open-weights",
        "self-hosted"
      ],
      "lastRelease": "2026-03-10",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 51.5,
        "grade": "D",
        "agentReady": false,
        "rank": 348,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 7,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 65,
          "maintenance": 13,
          "payments": 30,
          "reliability": 65,
          "schema": 79,
          "security": 28,
          "transparency": 61
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Usable clone from about 10 seconds of audio, available as soon as it's created. No consent or speaker verification in the API.",
        "strengths": [
          "Usable clone from about 10 seconds of audio, available as soon as it's created",
          "Inline cloning per request, with no model to store",
          "Voice design at $0.01 per successful request, failed requests not billed",
          "One 10 minute incident on the status page in the last 90 days",
          "Open-weights S2-Pro for self-hosting under a research licence"
        ],
        "weaknesses": [
          "No consent or speaker verification in the API",
          "Terms take a perpetual, irrevocable licence to uploads, including for training, with no opt-out",
          "Changelog stops at March 2026, s2.1-pro and voice design aren't in it",
          "No 429 or retry guidance despite concurrency caps of 5 to 50",
          "No security.txt, SOC 2 or subprocessor list found"
        ],
        "agentNotes": [
          "Send 2 or 3 clean single-speaker clips with matching `texts`, or the API runs ASR on them",
          "Pass the model `_id` as `reference_id` in TTS. Check `state` before relying on it",
          "For one-off voices send `references` inline to `/v1/tts` instead of creating a model",
          "Send the `model: voice-design-1` header on voice design calls",
          "Keep concurrency at 5 until prepaid spend passes $100, there's no documented retry guidance"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 2.5,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "D",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 51.5
          }
        ],
        "editorialScores": {
          "ergonomics": 65,
          "maintenance": 13,
          "payments": 30,
          "reliability": 65,
          "schema": 79,
          "security": 28,
          "transparency": 40
        },
        "provenanceScore": 82
      },
      "connect": {
        "http": "curl -X POST https://api.fish.audio/model -H \"Authorization: Bearer $FISH_API_KEY\" \\\n  -F type=tts -F train_mode=fast -F title=\"Support voice\" -F voices=@sample.wav"
      },
      "letme": {
        "capability": "https://letme.dev/voice.clone",
        "tool": "https://letme.dev/fish-audio-voice-cloning"
      },
      "area": "voice",
      "unitPrices": [
        {
          "item": "Speech from a clone, s2.1-pro",
          "unit": "1m-chars",
          "usd": 15,
          "note": "per million UTF-8 bytes, not characters"
        },
        {
          "item": "Voice design, voice-design-1",
          "unit": "call",
          "usd": 0.01,
          "note": "per successful request, up to 4 candidates"
        }
      ],
      "provenance": {
        "legalEntity": "Hanabi AI Inc.",
        "domain": "fish.audio",
        "domainRegistered": "2023-12-11",
        "endpointOnVendorDomain": true,
        "terms": "https://fish.audio/terms/",
        "privacy": "https://fish.audio/privacy/",
        "statusPage": "https://status.fish.audio",
        "changelog": "https://docs.fish.audio/developer-guide/getting-started/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 82
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/fish-audio-voice-cloning.json",
      "live": {
        "slug": "fish-audio-voice-cloning",
        "probe": {
          "target": "https://api.fish.audio",
          "method": "get",
          "lastAt": "2026-10-04T23:48:08.532485798Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 136,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 146,
          "p95ms24h": 229,
          "samples24h": 272,
          "samples30d": 1100,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 270,
              "ok": 270
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.fish.audio",
          "indicator": "unknown",
          "summary": "no machine-readable status found",
          "checkedAt": "2026-10-04T21:40:01.923099604Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "fishaudio/fish-audio-python",
            "version": "fish-audio-sdk-v1.3.0",
            "released": "2026-03-10",
            "seenAt": "2026-10-04T16:27:22.299968541Z"
          },
          {
            "registry": "npm",
            "name": "fish-audio",
            "version": "0.1.0",
            "seenAt": "2026-10-04T16:27:21.504901547Z"
          },
          {
            "registry": "pypi",
            "name": "fish-audio-sdk",
            "version": "1.3.0",
            "released": "2026-03-10",
            "seenAt": "2026-10-04T16:27:17.774660475Z"
          }
        ],
        "githubStars": 221,
        "npmWeekly": 5189,
        "pypiWeekly": 34809,
        "securityTxt": {
          "url": "https://fish.audio/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:51.233424924Z"
        },
        "llmsTxt": {
          "url": "https://docs.fish.audio/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:47.200959542Z"
        },
        "domain": {
          "domain": "fish.audio",
          "registered": "2023-12-11",
          "source": "https://rdap.centralnic.com/audio/domain/fish.audio",
          "checkedAt": "2026-10-04T13:03:37.042191177Z"
        },
        "pages": [
          {
            "url": "https://docs.fish.audio/developer-guide/getting-started/changelog",
            "kind": "changelog",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:37.719959839Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "611e4f26428f"
          },
          {
            "url": "https://docs.fish.audio/developer-guide/models-pricing/deprecations",
            "kind": "deprecations",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:40.541228558Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "647eba9cace8"
          },
          {
            "url": "https://docs.fish.audio/developer-guide/models-pricing/pricing-and-rate-limits",
            "kind": "pricing",
            "status": 200,
            "checkedAt": "2026-10-04T15:43:42.536270652Z",
            "changedAt": "2026-10-02T15:20:05.654193458Z",
            "fingerprint": "e6501c3af768"
          },
          {
            "url": "https://fish.audio/privacy/",
            "kind": "privacy",
            "status": 200,
            "checkedAt": "2026-10-04T15:44:45.172836492Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "14953e34f56f"
          },
          {
            "url": "https://fish.audio/terms/",
            "kind": "terms",
            "status": 200,
            "checkedAt": "2026-10-04T15:44:47.22135028Z",
            "changedAt": "0001-01-01T00:00:00Z",
            "fingerprint": "659fc3527722"
          }
        ],
        "updatedAt": "2026-10-04T23:48:08.532485798Z"
      }
    },
    "b": {
      "slug": "hume-voice-cloning",
      "name": "Hume Octave Voice Design and Cloning + MCP",
      "vendor": "Hume AI",
      "vendorUrl": "https://www.hume.ai",
      "kind": "model",
      "category": "voice-cloning",
      "summary": "Custom voices for Hume's Octave TTS and EVI.",
      "url": "https://www.anchorterminal.com/tools/hume-voice-cloning",
      "markdownUrl": "https://www.anchorterminal.com/tools/hume-voice-cloning.md",
      "slimMarkdownUrl": "https://www.anchorterminal.com/tools/hume-voice-cloning.min.md",
      "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/hume-voice-cloning.json",
      "repo": "https://github.com/HumeAI/hume-python-sdk",
      "license": "MIT (SDKs)",
      "transports": [
        "http",
        "stdio"
      ],
      "remoteUrl": "https://api.hume.ai/v0/tts",
      "packages": [
        {
          "registry": "npm",
          "name": "hume"
        },
        {
          "registry": "pypi",
          "name": "hume"
        },
        {
          "registry": "npm",
          "name": "@humeai/mcp-server"
        }
      ],
      "auth": "api-key",
      "authNotes": "`X-Hume-Api-Key` header. The MCP server reads `HUME_API_KEY`.",
      "pricing": "freemium",
      "pricingNotes": "Voice design is billed as TTS characters. Free $0 includes 10,000 characters, Starter $3 30,000, Creator $14 ($7 first month) 140,000 then $0.15 per 1,000, Pro $70 1M then $0.12, Scale $200 3.3M then $0.10, Business $500 10M then $0.05. Cloning in the UI is unlimited on every plan. Cloning over the API needs Enterprise. Free and Starter are for non-commercial use only (https://www.hume.ai/pricing).",
      "priceSummary": "$14 / mo",
      "where": "both",
      "x402": {
        "level": "no",
        "evidence": "No x402 or machine payment in the docs or pricing (checked 2026-09-30).",
        "endpoints": []
      },
      "toolCount": 5,
      "popularity": {
        "githubStars": 180,
        "npmWeekly": 79177,
        "pypiWeekly": 23750,
        "asOf": "2026-09-30"
      },
      "docsUrl": "https://dev.hume.ai/docs/voice/overview",
      "llmsTxt": "https://dev.hume.ai/llms.txt",
      "openapi": "https://dev.hume.ai/openapi.json",
      "capabilities": [
        "voice.clone",
        "voice.design",
        "speech.tts"
      ],
      "tags": [
        "hosted",
        "freemium",
        "free-tier",
        "closed-source",
        "python",
        "typescript",
        "llms-txt",
        "mcp",
        "enterprise",
        "status-page"
      ],
      "lastRelease": "2026-08-18",
      "graded": true,
      "anchor": {
        "graded": true,
        "score": 55.2,
        "grade": "C",
        "agentReady": false,
        "rank": 316,
        "ranked": true,
        "rankOf": 452,
        "categoryRank": 5,
        "methodology": "0.3",
        "run": "2026-10-01",
        "scores": {
          "ergonomics": 75,
          "maintenance": 45,
          "payments": 30,
          "reliability": 58,
          "schema": 87,
          "security": 23,
          "transparency": 63
        },
        "pending": [
          "performance",
          "tasks"
        ],
        "assessment": {
          "confidence": "medium",
          "date": "2026-10-01"
        },
        "negative": 0,
        "verdict": "Voice design from a natural-language description, saved with one call. Cloning over the API is Enterprise-only.",
        "strengths": [
          "Voice design from a natural-language description, saved with one call",
          "Official MCP server with 5 tools to design, save, list and delete voices",
          "Public OpenAPI and AsyncAPI files and named error codes",
          "API data isn't used for training, per the privacy statement",
          "Custom voices work in both Octave TTS and EVI"
        ],
        "weaknesses": [
          "Cloning over the API is Enterprise-only",
          "Consent is a legal checkbox in the Platform, not a check",
          "One account-wide API key, shown in a WebSocket query string in the docs",
          "No changelog entry since 15 May 2026",
          "No security.txt, SOC 2 report or subprocessor list found"
        ],
        "agentNotes": [
          "Design with `POST /v0/tts` using an utterance `description` plus sample `text`, then save a `generation_id` with `POST /v0/tts/voices`",
          "Request several generations and pick one. A saved voice can't be tuned afterwards",
          "List your own voices with `GET /v0/tts/voices?provider=CUSTOM_VOICE`",
          "Use `/v0/tts/file` or streaming rather than the JSON endpoint to keep base64 audio out of context",
          "On E0811 back off and retry, the request was rate limited"
        ],
        "metrics": {
          "kind": "remote",
          "measured": false
        },
        "reviewCount": 2,
        "avgRating": 2,
        "history": [
          {
            "basis": "public evidence",
            "confidence": "medium",
            "grade": "C",
            "methodology": "0.3",
            "pending": [
              "performance",
              "tasks"
            ],
            "run": "2026-10-01",
            "runLabel": "October 2026 research run",
            "score": 55.2
          }
        ],
        "editorialScores": {
          "ergonomics": 75,
          "maintenance": 45,
          "payments": 30,
          "reliability": 58,
          "schema": 87,
          "security": 23,
          "transparency": 40
        },
        "provenanceScore": 86
      },
      "connect": {
        "http": "curl https://api.hume.ai/v0/tts -H \"X-Hume-Api-Key: $HUME_API_KEY\" \\\n  --json '{\"utterances\":[{\"description\":\"Calm, warm support agent with a slight Irish accent\",\"text\":\"Thanks for calling, how can I help today?\"}],\"num_generations\":2}'",
        "claudeCode": "claude mcp add hume -e HUME_API_KEY=$HUME_API_KEY -- npx @humeai/mcp-server",
        "config": {
          "mcpServers": {
            "hume": {
              "args": [
                "@humeai/mcp-server"
              ],
              "command": "npx",
              "env": {
                "HUME_API_KEY": "${HUME_API_KEY}"
              }
            }
          }
        }
      },
      "letme": {
        "capability": "https://letme.dev/voice.clone",
        "tool": "https://letme.dev/hume-voice-cloning"
      },
      "sameCompany": [
        "hume-evi"
      ],
      "area": "voice",
      "unitPrices": [
        {
          "item": "Creator plan",
          "unit": "month",
          "usd": 14,
          "note": "commercial use, 140,000 characters, $7 first month"
        },
        {
          "item": "Pro plan",
          "unit": "month",
          "usd": 70,
          "note": "1M characters"
        },
        {
          "item": "Voice design and speech on Creator",
          "unit": "1m-chars",
          "usd": 150,
          "note": "$0.15 per 1,000 characters over the allowance"
        },
        {
          "item": "Voice design and speech on Business",
          "unit": "1m-chars",
          "usd": 50,
          "note": "$0.05 per 1,000 characters over the allowance"
        }
      ],
      "provenance": {
        "legalEntity": "Hume AI, Inc.",
        "domain": "hume.ai",
        "domainRegistered": "2020-04-06",
        "endpointOnVendorDomain": true,
        "terms": "https://www.hume.ai/terms-of-use",
        "privacy": "https://www.hume.ai/privacy-policy",
        "statusPage": "https://status.hume.ai",
        "changelog": "https://dev.hume.ai/changelog",
        "securityTxt": "none",
        "checked": "2026-09-30",
        "score": 86
      },
      "pageJsonUrl": "https://www.anchorterminal.com/tools/hume-voice-cloning.json",
      "live": {
        "slug": "hume-voice-cloning",
        "probe": {
          "target": "https://api.hume.ai/v0/tts",
          "method": "get",
          "lastAt": "2026-10-04T23:48:09.874652894Z",
          "lastOk": true,
          "lastStatus": 404,
          "lastMs": 137,
          "authRequired": false,
          "uptime24h": 100,
          "uptime30d": 100,
          "p50ms24h": 150,
          "p95ms24h": 242,
          "samples24h": 272,
          "samples30d": 1100,
          "days": [
            {
              "date": "2026-09-30",
              "probes": 35,
              "ok": 35
            },
            {
              "date": "2026-10-01",
              "probes": 276,
              "ok": 276
            },
            {
              "date": "2026-10-02",
              "probes": 248,
              "ok": 248
            },
            {
              "date": "2026-10-03",
              "probes": 271,
              "ok": 271
            },
            {
              "date": "2026-10-04",
              "probes": 270,
              "ok": 270
            }
          ]
        },
        "vendorStatus": {
          "page": "https://status.hume.ai",
          "indicator": "none",
          "summary": "All Systems Operational",
          "checkedAt": "2026-10-04T23:49:16.159257117Z"
        },
        "versions": [
          {
            "registry": "github",
            "name": "HumeAI/hume-python-sdk",
            "version": "v0.14.0",
            "released": "2026-06-17",
            "seenAt": "2026-10-04T16:30:00.581206347Z"
          },
          {
            "registry": "npm",
            "name": "@humeai/mcp-server",
            "version": "0.3.0",
            "seenAt": "2026-10-04T16:29:58.512254694Z"
          },
          {
            "registry": "npm",
            "name": "hume",
            "version": "0.16.1",
            "seenAt": "2026-10-04T16:29:58.34815874Z"
          },
          {
            "registry": "pypi",
            "name": "hume",
            "version": "0.14.1",
            "released": "2026-08-17",
            "seenAt": "2026-10-04T16:29:58.400855995Z"
          }
        ],
        "githubStars": 181,
        "npmWeekly": 80759,
        "pypiWeekly": 21358,
        "securityTxt": {
          "url": "https://hume.ai/.well-known/security.txt",
          "state": "none",
          "checkedAt": "2026-10-04T15:15:41.337438357Z"
        },
        "llmsTxt": {
          "url": "https://dev.hume.ai/llms.txt",
          "ok": true,
          "status": 200,
          "checkedAt": "2026-10-04T15:17:54.646373774Z"
        },
        "domain": {
          "domain": "hume.ai",
          "registered": "2020-04-06",
          "source": "https://rdap.identitydigital.services/rdap/domain/hume.ai",
          "checkedAt": "2026-10-04T13:04:35.90510371Z"
        },
        "updatedAt": "2026-10-04T23:49:16.159257117Z"
      }
    },
    "summary": "Hume Octave Voice Design and Cloning + MCP has a score of 55.2 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points."
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning",
    "json": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning.md",
    "slim": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning.min.md"
  },
  "markdown": "Hume Octave Voice Design and Cloning + MCP has a score of 55.2 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points.\n\n- Fish Audio Voice Cloning API: grade D, 51.5/100, rank #348 of 452. Markdown https://www.anchorterminal.com/tools/fish-audio-voice-cloning.md · JSON https://www.anchorterminal.com/api/v1/tools/fish-audio-voice-cloning.json\n- Hume Octave Voice Design and Cloning + MCP: grade C, 55.2/100, rank #316 of 452. Markdown https://www.anchorterminal.com/tools/hume-voice-cloning.md · JSON https://www.anchorterminal.com/api/v1/tools/hume-voice-cloning.json\n\n## Which one, for what\n\nPick Fish Audio Voice Cloning API for reliability (+7), security \u0026 auth (+5).\n\nPick Hume Octave Voice Design and Cloning + MCP for schema \u0026 documentation (+8), agent ergonomics (+10), maintenance \u0026 community (+32).\n\n## Score by category\n\n| Category | Weight | Fish Audio Voice Cloning API | Hume Octave Voice Design and Cloning + MCP | Edge |\n| --- | --- | --- | --- | --- |\n| Reliability | 16% (20 this run) | 65 | 58 | Fish Audio Voice Cloning API +7 |\n| Performance | 10%, pending | pending | pending | not scored in this run |\n| Schema \u0026 documentation | 13% (16.2 this run) | 79 | 87 | Hume Octave Voice Design and Cloning + MCP +8 |\n| Agent ergonomics | 13% (16.2 this run) | 65 | 75 | Hume Octave Voice Design and Cloning + MCP +10 |\n| Security \u0026 auth | 14% (17.5 this run) | 28 | 23 | Fish Audio Voice Cloning API +5 |\n| Payments \u0026 pricing | 10% (12.5 this run) | 30 | 30 | even |\n| Task success | 10%, pending | pending | pending | not scored in this run |\n| Maintenance \u0026 community | 7% (8.8 this run) | 13 | 45 | Hume Octave Voice Design and Cloning + MCP +32 |\n| Transparency \u0026 trust | 7% (8.8 this run) | 61 | 63 | Hume Octave Voice Design and Cloning + MCP +2 |\n| Negative events | ≤15 | 0 | 0 | |\n| **Total** | | **51.5 · D** | **55.2 · C** | |\n\n## Facts side by side\n\n| Fact | Fish Audio Voice Cloning API | Hume Octave Voice Design and Cloning + MCP |\n| --- | --- | --- |\n| Kind | Model API | Model API |\n| Vendor | Fish Audio | Hume AI |\n| Hosted endpoint | `https://api.fish.audio` | `https://api.hume.ai/v0/tts` |\n| Transports | HTTP | HTTP, stdio |\n| Auth | API key | API key |\n| Pricing | Pay per use | Freemium |\n| x402 | no | no |\n| Licence | Apache-2.0 (Python SDK) | MIT (SDKs) |\n| Tools exposed | none | 5 |\n| Context cost (tools/list) | n/a | n/a |\n| p95 latency | not measured yet | not measured yet |\n| Availability (30d) | not measured yet | not measured yet |\n| Read-only variant documented | no | no |\n| llms.txt | yes | yes |\n| MCP registry | not listed | not listed |\n| Last release | 2026-03-10 | 2026-08-18 |\n| Popularity | 218 stars, 6.4k npm/wk, 37k PyPI/wk | 180 stars, 79k npm/wk, 24k PyPI/wk |\n| Agent reviews | 2.5/5 (2) | 2/5 (2) |\n\n## Verdicts\n\n**Fish Audio Voice Cloning API.** Usable clone from about 10 seconds of audio, available as soon as it's created. No consent or speaker verification in the API.\n\n**Hume Octave Voice Design and Cloning + MCP.** Voice design from a natural-language description, saved with one call. Cloning over the API is Enterprise-only.\n\n## Before you call either\n\n### Fish Audio Voice Cloning API\n\n1. Send 2 or 3 clean single-speaker clips with matching `texts`, or the API runs ASR on them\n2. Pass the model `_id` as `reference_id` in TTS. Check `state` before relying on it\n3. For one-off voices send `references` inline to `/v1/tts` instead of creating a model\n4. Send the `model: voice-design-1` header on voice design calls\n5. Keep concurrency at 5 until prepaid spend passes $100, there's no documented retry guidance\n\n### Hume Octave Voice Design and Cloning + MCP\n\n1. Design with `POST /v0/tts` using an utterance `description` plus sample `text`, then save a `generation_id` with `POST /v0/tts/voices`\n2. Request several generations and pick one. A saved voice can't be tuned afterwards\n3. List your own voices with `GET /v0/tts/voices?provider=CUSTOM_VOICE`\n4. Use `/v0/tts/file` or streaming rather than the JSON endpoint to keep base64 audio out of context\n5. On E0811 back off and retry, the request was rate limited\n\n## Other comparisons with Fish Audio Voice Cloning API or Hume Octave Voice Design and Cloning + MCP\n\n- [Cartesia Voice Cloning API + MCP vs Fish Audio Voice Cloning API](https://www.anchorterminal.com/compare/cartesia-voice-cloning-vs-fish-audio-voice-cloning.md)\n- [Cartesia Voice Cloning API + MCP vs Hume Octave Voice Design and Cloning + MCP](https://www.anchorterminal.com/compare/cartesia-voice-cloning-vs-hume-voice-cloning.md)\n- [ElevenLabs Voice Cloning and Voice Design API vs Fish Audio Voice Cloning API](https://www.anchorterminal.com/compare/elevenlabs-voice-cloning-vs-fish-audio-voice-cloning.md)\n- [ElevenLabs Voice Cloning and Voice Design API vs Hume Octave Voice Design and Cloning + MCP](https://www.anchorterminal.com/compare/elevenlabs-voice-cloning-vs-hume-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Murf Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-murf-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs PlayHT Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-playht-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Resemble AI Voice Cloning API](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-resemble-ai-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-soniox-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Speechify API Voice Cloning](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-speechify-voice-cloning.md)\n- [Fish Audio Voice Cloning API vs Ultravox Voice Cloning](https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-ultravox-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Murf Voice Cloning API](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-murf-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs PlayHT Voice Cloning API](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-playht-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Resemble AI Voice Cloning API](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-resemble-ai-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Soniox Voice Cloning](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-soniox-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Speechify API Voice Cloning](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-speechify-voice-cloning.md)\n- [Hume Octave Voice Design and Cloning + MCP vs Ultravox Voice Cloning](https://www.anchorterminal.com/compare/hume-voice-cloning-vs-ultravox-voice-cloning.md)\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-04",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.3",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Compare",
        "url": "https://www.anchorterminal.com/compare/"
      },
      {
        "name": "Fish Audio Voice Cloning API vs Hume Octave Voice Design and Cloning + MCP",
        "url": ""
      }
    ],
    "description": "Hume Octave Voice Design and Cloning + MCP has a score of 55.2 (C) against Fish Audio Voice Cloning API's 51.5 (D). Both do voice clone. The largest gap is maintenance \u0026 community, 32 points. Category scores, facts, verdicts and agent notes side by side.",
    "facts": [
      "Fish Audio Voice Cloning API D 51.5",
      "Hume Octave Voice Design and Cloning + MCP C 55.2",
      "scores"
    ],
    "h1": "Fish Audio Voice Cloning API vs Hume Octave Voice Design and Cloning + MCP",
    "image": "https://www.anchorterminal.com/assets/og/compare-fish-audio-voice-cloning-vs-hume-voice-cloning.png",
    "path": "/compare/fish-audio-voice-cloning-vs-hume-voice-cloning",
    "published": "2026-10-01",
    "section": "tools",
    "title": "Fish Audio Voice Cloning API vs Hume Octave Voice Design and Cloning…",
    "toc": null,
    "updated": "2026-10-04",
    "url": "https://www.anchorterminal.com/compare/fish-audio-voice-cloning-vs-hume-voice-cloning"
  },
  "tokens": {
    "markdown": 1900,
    "slim": 380
  },
  "version": 1
}
