{
  "data": {
    "category": {
      "area": "models",
      "capabilities": [
        "embed.text",
        "embed.multimodal",
        "embed.code",
        "embed.multilingual",
        "rerank"
      ],
      "description": "Models that turn text, images and code into vectors for search, and rerankers that reorder search results by relevance. Compared on retrieval quality, dimensions and context length, multilingual support and price per million tokens.",
      "indexed": [
        {
          "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/elopstudio-ai-wave.json",
          "kind": "mcp",
          "name": "AI Wave",
          "slug": "elopstudio-ai-wave",
          "url": "https://www.anchorterminal.com/tools/elopstudio-ai-wave"
        },
        {
          "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/apple-rag-mcp-server.json",
          "kind": "mcp",
          "name": "apple-rag.com MCP server",
          "slug": "apple-rag-mcp-server",
          "url": "https://www.anchorterminal.com/tools/apple-rag-mcp-server"
        },
        {
          "jsonUrl": "https://www.anchorterminal.com/api/v1/tools/voxell-forge.json",
          "kind": "mcp",
          "name": "forge",
          "slug": "voxell-forge",
          "url": "https://www.anchorterminal.com/tools/voxell-forge"
        }
      ],
      "indexedCount": 3,
      "json": "https://www.anchorterminal.com/categories/embeddings.json",
      "name": "Embeddings \u0026 rerankers",
      "slug": "embeddings",
      "test": "The same corpus and queries embedded with each model, then the top results reranked. We check retrieval quality against labelled answers, latency and the cost per million tokens.",
      "title": "Embedding and reranking APIs for AI agents",
      "toolCount": 10,
      "tools": [
        "amazon-nova-embeddings",
        "openai-embeddings",
        "cohere-embed",
        "gemini-embedding",
        "jina-embeddings",
        "nvidia-nemo-retriever",
        "voyage-ai",
        "mistral-embeddings",
        "nomic-embed",
        "zeroentropy"
      ],
      "url": "https://www.anchorterminal.com/categories/embeddings"
    },
    "faq": [
      {
        "answer": "Amazon Nova Multimodal Embeddings has the highest benchmark score of the 10 ranked embedding and reranking APIs, 75 (BB). OpenAI embeddings is second with 73.2 (BB).",
        "question": "What are the highest-rated embedding and reranking APIs for AI agents?"
      },
      {
        "answer": "4 of the 10 ranked here grade BB or better, the bar for agent-ready on the Anchor benchmark.",
        "question": "How many embedding and reranking APIs are agent-ready?"
      },
      {
        "answer": "None of the ranked listings here accepts x402 for its main call yet.",
        "question": "Which embedding and reranking APIs accept x402 payments?"
      },
      {
        "answer": "By the Anchor benchmark score out of 100, a weighted mean of the scored categories minus deductions for negative events, from public evidence re-checked as vendors change. Listings cannot pay for a place. The latest assessment behind this page is from 8 October 2026.",
        "question": "How is this list ranked?"
      }
    ],
    "howToChoose": [
      {
        "label": "Vector dimensions and size",
        "detail": "Check the vector dimensions and whether they can be shortened, because the index stores and scans every dimension for every query."
      },
      {
        "label": "Query and document purposes",
        "detail": "Check whether queries and documents can be embedded with different purpose settings, since the setting can change which results rank highest."
      },
      {
        "label": "Languages and input types",
        "detail": "Confirm which languages and modalities the model supports, including images and code, because retrieval quality can fall outside those."
      },
      {
        "label": "Maximum input length",
        "detail": "Check the maximum input length and whether long documents are segmented for you, because a truncated chunk silently loses the text that answers a query."
      }
    ],
    "picks": [
      {
        "also": {
          "name": "OpenAI embeddings",
          "slug": "openai-embeddings",
          "why": "BB, 73.2/100"
        },
        "name": "Amazon Nova Multimodal Embeddings",
        "need": "Highest score overall",
        "slug": "amazon-nova-embeddings",
        "why": "BB, 75/100 on the benchmark"
      },
      {
        "name": "Cohere Embed and Rerank",
        "need": "Schema \u0026 documentation",
        "slug": "cohere-embed",
        "why": "92/100 on schema \u0026 documentation, against 76 for the overall leader"
      },
      {
        "name": "Voyage AI embeddings and rerankers",
        "need": "Agent ergonomics",
        "slug": "voyage-ai",
        "why": "98/100 on agent ergonomics, against 78 for the overall leader"
      },
      {
        "name": "OpenAI embeddings",
        "need": "Security \u0026 auth",
        "slug": "openai-embeddings",
        "why": "95/100 on security \u0026 auth, against 91 for the overall leader"
      },
      {
        "name": "Cohere Embed and Rerank",
        "need": "Maintenance \u0026 community",
        "slug": "cohere-embed",
        "why": "90/100 on maintenance \u0026 community, against 50 for the overall leader"
      },
      {
        "name": "OpenAI embeddings",
        "need": "Transparency \u0026 trust",
        "slug": "openai-embeddings",
        "why": "85/100 on transparency \u0026 trust, against 79 for the overall leader"
      },
      {
        "name": "Jina Embeddings and Reranker",
        "need": "A hosted MCP endpoint",
        "slug": "jina-embeddings",
        "why": "remote MCP server, nothing to install"
      },
      {
        "name": "NVIDIA NeMo Retriever Embedding and Reranking NIMs",
        "need": "Self-hosting under an open licence",
        "slug": "nvidia-nemo-retriever",
        "why": "self-hosted, Proprietary containers under the NVIDIA Software Licence Agreement and Product-Specific Terms for AI Products licence"
      }
    ],
    "ranked": 10,
    "shortlist": [
      {
        "bestFor": "Suited to mixed-media retrieval for teams already on AWS, especially video and audio archives processed through S3.",
        "grade": "BB",
        "name": "Amazon Nova Multimodal Embeddings",
        "position": 1,
        "price": "Pay per use",
        "score": 75,
        "slug": "amazon-nova-embeddings",
        "strengths": [
          "Text, images, document images, video and audio share one vector space, with four output sizes from 256 to 3072",
          "`embeddingPurpose` has nine documented values, with separate settings for indexing and for each retrieval type",
          "Published quotas of 2,000 requests a minute and 30 concurrent asynchronous jobs per Region"
        ],
        "url": "https://www.anchorterminal.com/tools/amazon-nova-embeddings",
        "verdict": "One model embeds text, images, document images, video and audio into a shared space, with nine documented purpose settings and published per-unit prices. It runs in US East (N. Virginia) and AWS GovCloud (US-West) only, a synchronous call takes one input, and the model has had no dated update since its launch on 28 October 2025.",
        "weaknesses": [
          "In-Region inference in us-east-1 and us-gov-west-1 only, with no cross-Region inference profile",
          "A synchronous request embeds one item, with at most 8,192 characters of inline text or 30 seconds of audio or video",
          "The Bedrock model card marks Invoke as unsupported while the Nova guide documents `InvokeModel` for synchronous calls"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "An agent already on OpenAI that needs cheap general-purpose text retrieval with a small index.",
        "grade": "BB",
        "name": "OpenAI embeddings",
        "position": 2,
        "price": "Pay per use",
        "score": 73.2,
        "slug": "openai-embeddings",
        "strengths": [
          "text-embedding-3-small at $0.02 per million tokens, $0.01 through the Batch API",
          "Restricted project keys are set per endpoint, so an agent's key can be cut down to read and model calls",
          "Up to 2,048 inputs and 300,000 tokens in one request"
        ],
        "url": "https://www.anchorterminal.com/tools/openai-embeddings",
        "verdict": "text-embedding-3-small at $0.02 per million tokens, $0.01 through the Batch API. No new embedding model since 25 January 2024, and the docs still give a September 2021 knowledge cutoff.",
        "weaknesses": [
          "No new embedding model since 25 January 2024, and the docs still give a September 2021 knowledge cutoff",
          "Text only, 8,192 tokens an input, and no reranker",
          "Over-long inputs fail rather than being truncated, and output is float or base64 only"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Best when reranking is the job, or for long multilingual documents and image-heavy material where a 128K embedding context helps, with a cheaper Fast model for queries against a Pro index.",
        "grade": "BB",
        "name": "Cohere Embed and Rerank",
        "position": 3,
        "price": "$2 / 1k req",
        "score": 72.5,
        "slug": "cohere-embed",
        "strengths": [
          "Embed 5 Pro and Fast share one embedding space, so a Pro index answers Fast queries",
          "128K context with six output sizes and int8, binary and base64 output",
          "Rerank 4 Pro and Fast with 32K context, top_n and published per-search prices"
        ],
        "url": "https://www.anchorterminal.com/tools/cohere-embed",
        "verdict": "Embed 5 Pro and Fast share one embedding space with 128K context and compressed outputs, and embed and rerank prices are public. Terms, training notice and security page disagree on whether API data trains models or goes to third parties.",
        "weaknesses": [
          "Terms, training notice and security page disagree on whether API data trains models or goes to third parties",
          "A Google Cloud outage degraded embed and rerank for about four hours on 1 September 2026, and Embed 5 isn't yet a status component",
          "96 inputs a call, and input_type is required"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Multimodal corpora, especially video and audio, and for agents already on Google Cloud.",
        "grade": "BB",
        "name": "Gemini Embedding",
        "position": 4,
        "price": "Freemium",
        "score": 70.6,
        "slug": "gemini-embedding",
        "strengths": [
          "Text, images, video, audio and PDFs interleaved in one request and one vector space",
          "Any output size from 128 to 3072, with truncated vectors returned normalised",
          "Batch API at half the standard embedding price"
        ],
        "url": "https://www.anchorterminal.com/tools/gemini-embedding",
        "verdict": "Text, images, video, audio and PDFs interleaved in one request and one vector space. $0.20 per million text tokens, against $0.02 for OpenAI's small model.",
        "weaknesses": [
          "$0.20 per million text tokens, against $0.02 for OpenAI's small model",
          "8,192 input tokens and float output only",
          "No reranker on the Gemini API"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Reranking large candidate sets and multimodal corpora with audio or video.",
        "grade": "C",
        "name": "Jina Embeddings and Reranker",
        "position": 5,
        "price": "Freemium",
        "score": 61,
        "slug": "jina-embeddings",
        "strengths": [
          "jina-reranker-v3.5 (20 July 2026) with a 131,072-token window and no document cap",
          "v5-omni embeds text, images, audio, video and PDFs into one space",
          "OpenAPI 3.1 file with enums for model, task and embedding_type, and error responses from 400 to 504"
        ],
        "url": "https://www.anchorterminal.com/tools/jina-embeddings",
        "verdict": "jina-reranker-v3.5 (20 July 2026) with a 131,072-token window and no document cap. No price per token in any currency on the public pages.",
        "weaknesses": [
          "No price per token in any currency on the public pages",
          "One prepaid balance shared with Reader and Search, so a scraping job can drain the embedding budget",
          "26 automated incidents on the status feed from 15 September to 1 October 2026, and no status component for v5-omni or reranker v3.5"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Teams that already run NVIDIA GPUs and need embedding and reranking inside their own network, including page-image retrieval with the VL models.",
        "grade": "C",
        "name": "NVIDIA NeMo Retriever Embedding and Reranking NIMs",
        "position": 6,
        "price": "Freemium",
        "score": 61,
        "slug": "nvidia-nemo-retriever",
        "strengths": [
          "OpenAPI 3.1 files for both services, with enums for `input_type`, `modality`, `embedding_type` and `truncate` and no extra properties allowed",
          "`/v1/embeddings` follows the OpenAI shape, and a `-query` or `-passage` model suffix replaces `input_type` for OpenAI clients",
          "Output can be shrunk with `dimensions` from 128 to 2048 on three models, or with `int8`, `uint8`, `binary` and `ubinary` types"
        ],
        "url": "https://www.anchorterminal.com/tools/nvidia-nemo-retriever",
        "verdict": "Self-hosted containers with OpenAPI 3.1 files, typed request fields, five embedding output types and a dated end-of-life list. The API has no authentication or rate limiting of its own, production use needs an NVIDIA AI Enterprise licence at $4,500 a GPU a year, and the release notes carry no dates.",
        "weaknesses": [
          "The API has no authentication and no rate limiting. The security page leaves both to a proxy the deployer runs",
          "Production use needs an NVIDIA AI Enterprise licence, $4,500 a GPU a year or $1 a GPU-hour on cloud marketplaces, plus the GPU",
          "Release notes carry no dates, and environment variable names changed between 2.0 and 2.3 without a note in the 2.2 or 2.3 notes we read"
        ],
        "where": "local",
        "x402": "no"
      },
      {
        "bestFor": "Retrieval quality across domains, code and long documents, with a reranker from the same key.",
        "grade": "C",
        "name": "Voyage AI embeddings and rerankers",
        "position": 7,
        "price": "Freemium",
        "score": 58.8,
        "slug": "voyage-ai",
        "strengths": [
          "200 million free tokens per current model, then $0.02 to $0.12 per million",
          "Domain models for code, finance and law, a multimodal model and contextualised chunk embeddings",
          "Output in float, int8, uint8, binary or ubinary at 256 to 2048 dimensions, per request"
        ],
        "url": "https://www.anchorterminal.com/tools/voyage-ai",
        "verdict": "200 million free tokens per current model, then $0.02 to $0.12 per million. Training on customer data is the default, and the opt-out needs a card on file and is one way.",
        "weaknesses": [
          "Training on customer data is the default, and the opt-out needs a card on file and is one way",
          "No security.txt, and no status page linked or reachable",
          "Rate-limit tiers only begin once a payment method is added"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "EU data residency, Mistral-only stacks and code retrieval with small vectors.",
        "grade": "C",
        "name": "Mistral Embed and Codestral Embed",
        "position": 8,
        "price": "Freemium",
        "score": 57.9,
        "slug": "mistral-embeddings",
        "strengths": [
          "EU and US regional endpoints and a French legal entity",
          "codestral-embed with up to 3072 dimensions, first-n truncation and int8 or binary output",
          "OpenAPI document and llms.txt for the whole API"
        ],
        "url": "https://www.anchorterminal.com/tools/mistral-embeddings",
        "verdict": "EU and US regional endpoints and a French legal entity. Embedding API uptime of 94.36 per cent over 90 days on Mistral's status page, with incidents on 12 and 27 August 2026.",
        "weaknesses": [
          "Embedding API uptime of 94.36 per cent over 90 days on Mistral's status page, with incidents on 12 and 27 August 2026",
          "8k context on both models, and text or code only",
          "mistral-embed dates from December 2023 with fixed 1024-dimension float output, and nothing new since May 2025"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Teams that want a hosted endpoint for an open-weight model they can also run themselves, with the same vectors either way.",
        "grade": "D",
        "name": "Nomic Embed",
        "position": 9,
        "price": "Freemium",
        "score": 49.2,
        "slug": "nomic-embed",
        "strengths": [
          "Weights for nomic-embed-text-v1, v1.5, v2-moe, nomic-embed-code and nomic-embed-vision-v1.5 are Apache-2.0 on Hugging Face",
          "Public OpenAPI 3.1 document for the Atlas API (v0.57.0) with typed request and response schemas for both embedding endpoints",
          "nomic-embed-text-v1.5 accepts a dimensionality from 64 to 768, and inputs up to 8,192 tokens per text"
        ],
        "url": "https://www.anchorterminal.com/tools/nomic-embed",
        "verdict": "The text models have Apache-2.0 weights and a public OpenAPI 3.1 contract, so vectors made through the hosted endpoint can be reproduced locally. Nomic's current site and documentation index describe a construction-industry product, no rendered public page prices the endpoint, and no published terms or status component name it.",
        "weaknesses": [
          "docs.nomic.ai/llms.txt and www.nomic.ai now describe a product for architecture, engineering and construction firms, and the documentation index no longer lists the embedding pages",
          "No rendered public page states a price for the endpoint. The $1 per 10M tokens figure comes from the Atlas web app's script",
          "status.nomic.ai has no component for api-atlas.nomic.ai, and no SLA was found"
        ],
        "where": "hosted",
        "x402": "no"
      },
      {
        "bestFor": "Only as open weights for teams that can self-host a reranker or embedding model.",
        "grade": "F",
        "name": "ZeroEntropy zerank and zembed",
        "position": 10,
        "price": "Pay per use",
        "score": 13.7,
        "slug": "zeroentropy",
        "strengths": [
          "All four models now open weights under Apache 2.0 on Hugging Face",
          "A migration guide with self-hosting recipes for Baseten and Modal and named hosted alternatives",
          "42 days' notice before the API was retired"
        ],
        "url": "https://www.anchorterminal.com/tools/zeroentropy",
        "verdict": "All four models now open weights under Apache 2.0 on Hugging Face. The hosted API was discontinued after 4 September 2026 and signups closed on 24 July 2026.",
        "weaknesses": [
          "The hosted API was discontinued after 4 September 2026 and signups closed on 24 July 2026",
          "The docs and pricing page still advertise per-token API prices without mentioning the shutdown",
          "Nothing published on what happens to customer documents after the shutdown"
        ],
        "where": "hosted",
        "x402": "no"
      }
    ],
    "updated": "2026-10-08"
  },
  "kind": "anchor.page",
  "links": {
    "api": "https://www.anchorterminal.com/api/v1/index.json",
    "html": "https://www.anchorterminal.com/best/embeddings/",
    "json": "https://www.anchorterminal.com/best/embeddings/index.json",
    "llms": "https://www.anchorterminal.com/llms.txt",
    "markdown": "https://www.anchorterminal.com/best/embeddings/index.md",
    "slim": "https://www.anchorterminal.com/best/embeddings/index.min.md"
  },
  "markdown": "All 10 ranked embedding and reranking APIs on the Anchor benchmark, with a pick for each need and where each one falls short. Scores come from public evidence, re-checked as vendors change.\n\n- Ranked: 10 · agent-ready (BB or better): 4 · accept x402: 0 · hosted endpoints: 9\n- Full ranked table: https://www.anchorterminal.com/categories/embeddings.md\n- Head-to-head comparisons: https://www.anchorterminal.com/compare/embeddings/index.md (45)\n- Methodology: https://www.anchorterminal.com/benchmark/index.md\n\n## The shortlist\n\n| # | Tool | Grade | Score | Best for | Price | Where |\n| --- | --- | --- | --- | --- | --- | --- |\n| 1 | [Amazon Nova Multimodal Embeddings](https://www.anchorterminal.com/tools/amazon-nova-embeddings.md) | BB | 75 | Suited to mixed-media retrieval for teams already on AWS, especially video and audio archives processed through S3. | Pay per use | hosted |\n| 2 | [OpenAI embeddings](https://www.anchorterminal.com/tools/openai-embeddings.md) | BB | 73.2 | An agent already on OpenAI that needs cheap general-purpose text retrieval with a small index. | Pay per use | hosted |\n| 3 | [Cohere Embed and Rerank](https://www.anchorterminal.com/tools/cohere-embed.md) | BB | 72.5 | Best when reranking is the job, or for long multilingual documents and image-heavy material where a 128K embedding context helps, with a cheaper Fast model for queries against a Pro index. | $2 / 1k req | hosted |\n| 4 | [Gemini Embedding](https://www.anchorterminal.com/tools/gemini-embedding.md) | BB | 70.6 | Multimodal corpora, especially video and audio, and for agents already on Google Cloud. | Freemium | hosted |\n| 5 | [Jina Embeddings and Reranker](https://www.anchorterminal.com/tools/jina-embeddings.md) | C | 61 | Reranking large candidate sets and multimodal corpora with audio or video. | Freemium | hosted |\n| 6 | [NVIDIA NeMo Retriever Embedding and Reranking NIMs](https://www.anchorterminal.com/tools/nvidia-nemo-retriever.md) | C | 61 | Teams that already run NVIDIA GPUs and need embedding and reranking inside their own network, including page-image retrieval with the VL models. | Freemium | local |\n| 7 | [Voyage AI embeddings and rerankers](https://www.anchorterminal.com/tools/voyage-ai.md) | C | 58.8 | Retrieval quality across domains, code and long documents, with a reranker from the same key. | Freemium | hosted |\n| 8 | [Mistral Embed and Codestral Embed](https://www.anchorterminal.com/tools/mistral-embeddings.md) | C | 57.9 | EU data residency, Mistral-only stacks and code retrieval with small vectors. | Freemium | hosted |\n| 9 | [Nomic Embed](https://www.anchorterminal.com/tools/nomic-embed.md) | D | 49.2 | Teams that want a hosted endpoint for an open-weight model they can also run themselves, with the same vectors either way. | Freemium | hosted |\n| 10 | [ZeroEntropy zerank and zembed](https://www.anchorterminal.com/tools/zeroentropy.md) | F | 13.7 | Only as open weights for teams that can self-host a reranker or embedding model. | Pay per use | hosted |\n\n## Picks by need\n\n- Highest score overall: [Amazon Nova Multimodal Embeddings](https://www.anchorterminal.com/tools/amazon-nova-embeddings.md), BB, 75/100 on the benchmark. Also [OpenAI embeddings](https://www.anchorterminal.com/tools/openai-embeddings.md), BB, 73.2/100.\n- Schema \u0026 documentation: [Cohere Embed and Rerank](https://www.anchorterminal.com/tools/cohere-embed.md), 92/100 on schema \u0026 documentation, against 76 for the overall leader.\n- Agent ergonomics: [Voyage AI embeddings and rerankers](https://www.anchorterminal.com/tools/voyage-ai.md), 98/100 on agent ergonomics, against 78 for the overall leader.\n- Security \u0026 auth: [OpenAI embeddings](https://www.anchorterminal.com/tools/openai-embeddings.md), 95/100 on security \u0026 auth, against 91 for the overall leader.\n- Maintenance \u0026 community: [Cohere Embed and Rerank](https://www.anchorterminal.com/tools/cohere-embed.md), 90/100 on maintenance \u0026 community, against 50 for the overall leader.\n- Transparency \u0026 trust: [OpenAI embeddings](https://www.anchorterminal.com/tools/openai-embeddings.md), 85/100 on transparency \u0026 trust, against 79 for the overall leader.\n- A hosted MCP endpoint: [Jina Embeddings and Reranker](https://www.anchorterminal.com/tools/jina-embeddings.md), remote MCP server, nothing to install.\n- Self-hosting under an open licence: [NVIDIA NeMo Retriever Embedding and Reranking NIMs](https://www.anchorterminal.com/tools/nvidia-nemo-retriever.md), self-hosted, Proprietary containers under the NVIDIA Software Licence Agreement and Product-Specific Terms for AI Products licence.\n\n## How to choose\n\n- Vector dimensions and size: Check the vector dimensions and whether they can be shortened, because the index stores and scans every dimension for every query.\n- Query and document purposes: Check whether queries and documents can be embedded with different purpose settings, since the setting can change which results rank highest.\n- Languages and input types: Confirm which languages and modalities the model supports, including images and code, because retrieval quality can fall outside those.\n- Maximum input length: Check the maximum input length and whether long documents are segmented for you, because a truncated chunk silently loses the text that answers a query.\n\n- How the benchmark tests this category: The same corpus and queries embedded with each model, then the top results reranked. We check retrieval quality against labelled answers, latency and the cost per million tokens.\n\n## Each one in detail\n\n### 1. Amazon Nova Multimodal Embeddings, BB 75/100\n\nAmazon Nova Multimodal Embeddings is an AWS model on Amazon Bedrock that turns text, images, document images, video and audio into vectors in one space, at 256, 384, 1024 or 3072 dimensions, through synchronous and asynchronous calls.\n\n- Verdict: One model embeds text, images, document images, video and audio into a shared space, with nine documented purpose settings and published per-unit prices. It runs in US East (N. Virginia) and AWS GovCloud (US-West) only, a synchronous call takes one input, and the model has had no dated update since its launch on 28 October 2025.\n- Choose it for: Suited to mixed-media retrieval for teams already on AWS, especially video and audio archives processed through S3.\n- Strength: Text, images, document images, video and audio share one vector space, with four output sizes from 256 to 3072\n- Strength: `embeddingPurpose` has nine documented values, with separate settings for indexing and for each retrieval type\n- Strength: Published quotas of 2,000 requests a minute and 30 concurrent asynchronous jobs per Region\n- Weakness: In-Region inference in us-east-1 and us-gov-west-1 only, with no cross-Region inference profile\n- Weakness: A synchronous request embeds one item, with at most 8,192 characters of inline text or 30 seconds of audio or video\n- Weakness: The Bedrock model card marks Invoke as unsupported while the Nova guide documents `InvokeModel` for synchronous calls\n- Price: Pay per use · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/amazon-nova-embeddings.md\n\n### 2. OpenAI embeddings, BB 73.2/100\n\nOpenAI's text embedding API, with adjustable output dimensions for search and retrieval applications.\n\n- Verdict: text-embedding-3-small at $0.02 per million tokens, $0.01 through the Batch API. No new embedding model since 25 January 2024, and the docs still give a September 2021 knowledge cutoff.\n- Choose it for: An agent already on OpenAI that needs cheap general-purpose text retrieval with a small index.\n- Strength: text-embedding-3-small at $0.02 per million tokens, $0.01 through the Batch API\n- Strength: Restricted project keys are set per endpoint, so an agent's key can be cut down to read and model calls\n- Strength: Up to 2,048 inputs and 300,000 tokens in one request\n- Weakness: No new embedding model since 25 January 2024, and the docs still give a September 2021 knowledge cutoff\n- Weakness: Text only, 8,192 tokens an input, and no reranker\n- Weakness: Over-long inputs fail rather than being truncated, and output is float or base64 only\n- Price: Pay per use · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/openai-embeddings.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-openai-embeddings.md\n\n### 3. Cohere Embed and Rerank, BB 72.5/100\n\nCohere's Embed API turns text, images and mixed text-and-image inputs such as PDF pages into vectors with Embed 5 Pro and Fast, and its Rerank API reorders search results with Rerank 4.\n\n- Verdict: Embed 5 Pro and Fast share one embedding space with 128K context and compressed outputs, and embed and rerank prices are public. Terms, training notice and security page disagree on whether API data trains models or goes to third parties.\n- Choose it for: Best when reranking is the job, or for long multilingual documents and image-heavy material where a 128K embedding context helps, with a cheaper Fast model for queries against a Pro index.\n- Strength: Embed 5 Pro and Fast share one embedding space, so a Pro index answers Fast queries\n- Strength: 128K context with six output sizes and int8, binary and base64 output\n- Strength: Rerank 4 Pro and Fast with 32K context, top_n and published per-search prices\n- Weakness: Terms, training notice and security page disagree on whether API data trains models or goes to third parties\n- Weakness: A Google Cloud outage degraded embed and rerank for about four hours on 1 September 2026, and Embed 5 isn't yet a status component\n- Weakness: 96 inputs a call, and input_type is required\n- Price: $2 / 1k req · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/cohere-embed.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-cohere-embed.md\n\n### 4. Gemini Embedding, BB 70.6/100\n\ngemini-embedding-2, Google's multimodal embedding model, takes text, images, video, audio and PDFs into one 3072-dimension space (truncatable to 128) at 8,192 input tokens in 100+ languages.\n\n- Verdict: Text, images, video, audio and PDFs interleaved in one request and one vector space. $0.20 per million text tokens, against $0.02 for OpenAI's small model.\n- Choose it for: Multimodal corpora, especially video and audio, and for agents already on Google Cloud.\n- Strength: Text, images, video, audio and PDFs interleaved in one request and one vector space\n- Strength: Any output size from 128 to 3072, with truncated vectors returned normalised\n- Strength: Batch API at half the standard embedding price\n- Weakness: $0.20 per million text tokens, against $0.02 for OpenAI's small model\n- Weakness: 8,192 input tokens and float output only\n- Weakness: No reranker on the Gemini API\n- Price: Freemium · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/gemini-embedding.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-gemini-embedding.md\n\n### 5. Jina Embeddings and Reranker, C 61/100\n\njina-embeddings-v5 in text and omni (text, image, audio, video, PDF) variants at up to 32,768 tokens, plus the jina-reranker-v3.5 at 131,072 tokens a call.\n\n- Verdict: jina-reranker-v3.5 (20 July 2026) with a 131,072-token window and no document cap. No price per token in any currency on the public pages.\n- Choose it for: Reranking large candidate sets and multimodal corpora with audio or video.\n- Strength: jina-reranker-v3.5 (20 July 2026) with a 131,072-token window and no document cap\n- Strength: v5-omni embeds text, images, audio, video and PDFs into one space\n- Strength: OpenAPI 3.1 file with enums for model, task and embedding_type, and error responses from 400 to 504\n- Weakness: No price per token in any currency on the public pages\n- Weakness: One prepaid balance shared with Reader and Search, so a scraping job can drain the embedding budget\n- Weakness: 26 automated incidents on the status feed from 15 September to 1 October 2026, and no status component for v5-omni or reranker v3.5\n- Price: Freemium · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/jina-embeddings.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-jina-embeddings.md\n\n### 6. NVIDIA NeMo Retriever Embedding and Reranking NIMs, C 61/100\n\nNVIDIA's NeMo Retriever Embedding and Reranking NIMs are GPU containers that run text and image embedding models and rerankers behind a local REST API, with `/v1/embeddings` in the OpenAI shape and `/v1/ranking`.\n\n- Verdict: Self-hosted containers with OpenAPI 3.1 files, typed request fields, five embedding output types and a dated end-of-life list. The API has no authentication or rate limiting of its own, production use needs an NVIDIA AI Enterprise licence at $4,500 a GPU a year, and the release notes carry no dates.\n- Choose it for: Teams that already run NVIDIA GPUs and need embedding and reranking inside their own network, including page-image retrieval with the VL models.\n- Strength: OpenAPI 3.1 files for both services, with enums for `input_type`, `modality`, `embedding_type` and `truncate` and no extra properties allowed\n- Strength: `/v1/embeddings` follows the OpenAI shape, and a `-query` or `-passage` model suffix replaces `input_type` for OpenAI clients\n- Strength: Output can be shrunk with `dimensions` from 128 to 2048 on three models, or with `int8`, `uint8`, `binary` and `ubinary` types\n- Weakness: The API has no authentication and no rate limiting. The security page leaves both to a proxy the deployer runs\n- Weakness: Production use needs an NVIDIA AI Enterprise licence, $4,500 a GPU a year or $1 a GPU-hour on cloud marketplaces, plus the GPU\n- Weakness: Release notes carry no dates, and environment variable names changed between 2.0 and 2.3 without a note in the 2.2 or 2.3 notes we read\n- Price: Freemium · Auth: None · x402: no · Where: local\n- Full assessment: https://www.anchorterminal.com/tools/nvidia-nemo-retriever.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-nvidia-nemo-retriever.md\n\n### 7. Voyage AI embeddings and rerankers, C 58.8/100\n\nEmbedding and reranking models for text, code and multimodal retrieval from MongoDB-owned Voyage AI.\n\n- Verdict: 200 million free tokens per current model, then $0.02 to $0.12 per million. Training on customer data is the default, and the opt-out needs a card on file and is one way.\n- Choose it for: Retrieval quality across domains, code and long documents, with a reranker from the same key.\n- Strength: 200 million free tokens per current model, then $0.02 to $0.12 per million\n- Strength: Domain models for code, finance and law, a multimodal model and contextualised chunk embeddings\n- Strength: Output in float, int8, uint8, binary or ubinary at 256 to 2048 dimensions, per request\n- Weakness: Training on customer data is the default, and the opt-out needs a card on file and is one way\n- Weakness: No security.txt, and no status page linked or reachable\n- Weakness: Rate-limit tiers only begin once a payment method is added\n- Price: Freemium · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/voyage-ai.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-voyage-ai.md\n\n### 8. Mistral Embed and Codestral Embed, C 57.9/100\n\nMistral's API for generating text and code embeddings.\n\n- Verdict: EU and US regional endpoints and a French legal entity. Embedding API uptime of 94.36 per cent over 90 days on Mistral's status page, with incidents on 12 and 27 August 2026.\n- Choose it for: EU data residency, Mistral-only stacks and code retrieval with small vectors.\n- Strength: EU and US regional endpoints and a French legal entity\n- Strength: codestral-embed with up to 3072 dimensions, first-n truncation and int8 or binary output\n- Strength: OpenAPI document and llms.txt for the whole API\n- Weakness: Embedding API uptime of 94.36 per cent over 90 days on Mistral's status page, with incidents on 12 and 27 August 2026\n- Weakness: 8k context on both models, and text or code only\n- Weakness: mistral-embed dates from December 2023 with fixed 1024-dimension float output, and nothing new since May 2025\n- Price: Freemium · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/mistral-embeddings.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-mistral-embeddings.md\n\n### 9. Nomic Embed, D 49.2/100\n\nNomic's hosted embedding endpoints on the Atlas API turn text and images into vectors with the open-weight Nomic Embed models. Agents call them over HTTP with an API key or through the Python and TypeScript clients.\n\n- Verdict: The text models have Apache-2.0 weights and a public OpenAPI 3.1 contract, so vectors made through the hosted endpoint can be reproduced locally. Nomic's current site and documentation index describe a construction-industry product, no rendered public page prices the endpoint, and no published terms or status component name it.\n- Choose it for: Teams that want a hosted endpoint for an open-weight model they can also run themselves, with the same vectors either way.\n- Strength: Weights for nomic-embed-text-v1, v1.5, v2-moe, nomic-embed-code and nomic-embed-vision-v1.5 are Apache-2.0 on Hugging Face\n- Strength: Public OpenAPI 3.1 document for the Atlas API (v0.57.0) with typed request and response schemas for both embedding endpoints\n- Strength: nomic-embed-text-v1.5 accepts a dimensionality from 64 to 768, and inputs up to 8,192 tokens per text\n- Weakness: docs.nomic.ai/llms.txt and www.nomic.ai now describe a product for architecture, engineering and construction firms, and the documentation index no longer lists the embedding pages\n- Weakness: No rendered public page states a price for the endpoint. The $1 per 10M tokens figure comes from the Atlas web app's script\n- Weakness: status.nomic.ai has no component for api-atlas.nomic.ai, and no SLA was found\n- Price: Freemium · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/nomic-embed.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-nomic-embed.md\n\n### 10. ZeroEntropy zerank and zembed, F 13.7/100\n\nDiscontinued retrieval API acquired by Notion. Its embedding and reranking models remain available as open weights for self-hosting.\n\n- Verdict: All four models now open weights under Apache 2.0 on Hugging Face. The hosted API was discontinued after 4 September 2026 and signups closed on 24 July 2026.\n- Choose it for: Only as open weights for teams that can self-host a reranker or embedding model.\n- Strength: All four models now open weights under Apache 2.0 on Hugging Face\n- Strength: A migration guide with self-hosting recipes for Baseten and Modal and named hosted alternatives\n- Strength: 42 days' notice before the API was retired\n- Weakness: The hosted API was discontinued after 4 September 2026 and signups closed on 24 July 2026\n- Weakness: The docs and pricing page still advertise per-token API prices without mentioning the shutdown\n- Weakness: Nothing published on what happens to customer documents after the shutdown\n- Price: Pay per use · Auth: API key · x402: no · Where: hosted\n- Full assessment: https://www.anchorterminal.com/tools/zeroentropy.md\n- Against #1: https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-zeroentropy.md\n\n## Head to head\n\n- [Amazon Nova Multimodal Embeddings vs OpenAI embeddings](https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-openai-embeddings.md)\n- [Amazon Nova Multimodal Embeddings vs Cohere Embed and Rerank](https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-cohere-embed.md)\n- [Amazon Nova Multimodal Embeddings vs Gemini Embedding](https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-gemini-embedding.md)\n- [Amazon Nova Multimodal Embeddings vs Jina Embeddings and Reranker](https://www.anchorterminal.com/compare/amazon-nova-embeddings-vs-jina-embeddings.md)\n- [Cohere Embed and Rerank vs OpenAI embeddings](https://www.anchorterminal.com/compare/cohere-embed-vs-openai-embeddings.md)\n- [Gemini Embedding vs OpenAI embeddings](https://www.anchorterminal.com/compare/gemini-embedding-vs-openai-embeddings.md)\n- [Jina Embeddings and Reranker vs OpenAI embeddings](https://www.anchorterminal.com/compare/jina-embeddings-vs-openai-embeddings.md)\n- [Cohere Embed and Rerank vs Gemini Embedding](https://www.anchorterminal.com/compare/cohere-embed-vs-gemini-embedding.md)\n- [Cohere Embed and Rerank vs Jina Embeddings and Reranker](https://www.anchorterminal.com/compare/cohere-embed-vs-jina-embeddings.md)\n- [Gemini Embedding vs Jina Embeddings and Reranker](https://www.anchorterminal.com/compare/gemini-embedding-vs-jina-embeddings.md)\n\n## Questions\n\n### What are the highest-rated embedding and reranking APIs for AI agents?\n\nAmazon Nova Multimodal Embeddings has the highest benchmark score of the 10 ranked embedding and reranking APIs, 75 (BB). OpenAI embeddings is second with 73.2 (BB).\n\n### How many embedding and reranking APIs are agent-ready?\n\n4 of the 10 ranked here grade BB or better, the bar for agent-ready on the Anchor benchmark.\n\n### Which embedding and reranking APIs accept x402 payments?\n\nNone of the ranked listings here accepts x402 for its main call yet.\n\n### How is this list ranked?\n\nBy the Anchor benchmark score out of 100, a weighted mean of the scored categories minus deductions for negative events, from public evidence re-checked as vendors change. Listings cannot pay for a place. The latest assessment behind this page is from 8 October 2026.\n\n## How this list is made\n\nThe order is the Anchor benchmark score, the same number as on each listing. Each listing is graded from public evidence against the benchmark checklist, and the picks are worked out from those grades, prices and facts. No listing pays for its place, and paid audits or listing help never change a score.\n",
  "meta": {
    "attribution": "Anchor Terminal (https://www.anchorterminal.com)",
    "docs": "https://www.anchorterminal.com/docs/",
    "generatedAt": "2026-10-10",
    "license": "CC-BY-4.0",
    "method": "https://www.anchorterminal.com/benchmark/",
    "methodology": "0.4",
    "openapi": "https://www.anchorterminal.com/openapi.json",
    "preview": false,
    "run": "2026-10-01",
    "runLabel": "October 2026 research run"
  },
  "page": {
    "breadcrumbs": [
      {
        "name": "Home",
        "url": "https://www.anchorterminal.com/"
      },
      {
        "name": "Best of",
        "url": "https://www.anchorterminal.com/best/"
      },
      {
        "name": "Embeddings \u0026 rerankers",
        "url": ""
      }
    ],
    "description": "Amazon Nova Multimodal Embeddings (BB), OpenAI embeddings (BB) and Cohere Embed and Rerank (BB) lead the 10 ranked embedding and reranking APIs. Picks by need, strengths, weaknesses and prices from the Anchor benchmark.",
    "facts": [
      "Amazon Nova Multimodal Embeddings BB",
      "OpenAI embeddings BB",
      "Cohere Embed and Rerank BB"
    ],
    "h1": "Best embedding and reranking APIs for AI agents",
    "image": "https://www.anchorterminal.com/assets/og/best-embeddings.png",
    "path": "/best/embeddings/",
    "published": "",
    "section": "tools",
    "title": "Best embedding and reranking APIs for AI agents in 2026, ranked",
    "toc": null,
    "updated": "2026-10-08",
    "url": "https://www.anchorterminal.com/best/embeddings/"
  },
  "tokens": {
    "markdown": 5700,
    "slim": 1530
  },
  "version": 1
}
