{
  "alternatives": [
    {
      "auth": "api-key",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "endpoint": "https://transcribe.us-east-1.amazonaws.com",
      "grade": "BB",
      "json": "https://www.anchorterminal.com/api/v1/tools/amazon-transcribe.json",
      "kind": "model",
      "leadPriceUsd": 0.0024,
      "leadUnit": "audio-minute",
      "letme": "https://letme.dev/amazon-transcribe",
      "markdown": "https://www.anchorterminal.com/tools/amazon-transcribe.md",
      "name": "Amazon Transcribe",
      "priceSummary": "Pay per use",
      "pricing": "usage",
      "rank": 57,
      "sample": false,
      "score": 73.6,
      "slug": "amazon-transcribe",
      "url": "https://www.anchorterminal.com/tools/amazon-transcribe",
      "vendor": "Amazon Web Services",
      "where": "hosted",
      "x402": "no"
    },
    {
      "auth": "api-key",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages"
      ],
      "endpoint": "https://api.deepgram.com/v1",
      "grade": "BB",
      "json": "https://www.anchorterminal.com/api/v1/tools/deepgram-stt.json",
      "kind": "model",
      "leadPriceUsd": 0.002,
      "leadUnit": "audio-minute",
      "letme": "https://letme.dev/deepgram-stt",
      "markdown": "https://www.anchorterminal.com/tools/deepgram-stt.md",
      "name": "Deepgram Speech-to-Text (Nova-3, Flux)",
      "priceSummary": "Pay per use",
      "pricing": "usage",
      "rank": 94,
      "sample": false,
      "score": 70.6,
      "slug": "deepgram-stt",
      "url": "https://www.anchorterminal.com/tools/deepgram-stt",
      "vendor": "Deepgram",
      "where": "both",
      "x402": "no"
    },
    {
      "auth": "oauth",
      "capabilities": [
        "speech.stt",
        "speech.streaming",
        "speech.batch",
        "speech.diarisation",
        "speech.languages",
        "speech.translation"
      ],
      "endpoint": "https://speech.googleapis.com/v2",
      "grade": "BB",
      "json": "https://www.anchorterminal.com/api/v1/tools/google-speech-to-text.json",
      "kind": "model",
      "leadPriceUsd": 0.003,
      "leadUnit": "audio-minute",
      "letme": "https://letme.dev/google-speech-to-text",
      "markdown": "https://www.anchorterminal.com/tools/google-speech-to-text.md",
      "name": "Google Cloud Speech-to-Text",
      "priceSummary": "Freemium",
      "pricing": "freemium",
      "rank": 98,
      "sample": false,
      "score": 70.4,
      "slug": "google-speech-to-text",
      "url": "https://www.anchorterminal.com/tools/google-speech-to-text",
      "vendor": "Google Cloud",
      "where": "hosted",
      "x402": "no"
    }
  ],
  "calling": {
    "note": "letme picks today. Calling through letme (one key, payments) comes later; call the pick direct with the details above.",
    "open": false
  },
  "capability": "speech.stt",
  "directory": {
    "base": "https://www.anchorterminal.com",
    "build": "20261004T190737Z",
    "fetchedAt": "2026-10-04T21:21:02Z",
    "note": "Anchor Terminal grades the tools and publishes the grades; letme reads them from the same public files anyone can read. letme's terms are never an input to a grade.",
    "sameCompany": true
  },
  "filters": {
    "minGrade": "BB",
    "n": 10
  },
  "letme": "https://letme.dev/speech.stt?minGrade=BB&n=10",
  "live": false,
  "order": "in the rule's order",
  "pick": {
    "agentNotes": [
      "Use fast transcription (`transcriptions:transcribe`) for files under 5 hours and 500 MB, and batch for bulk jobs",
      "Pin `api-version=2025-10-15`. v3.0 and the v3.2 previews are retired",
      "On a 429, back off 1, 2, 4 then 4 minutes. It usually means autoscaling, not a quota"
    ],
    "auth": "mixed",
    "authNotes": "`Ocp-Apim-Subscription-Key` header with a Speech resource key, or a Microsoft Entra ID bearer token (Microsoft's recommended keyless option). Endpoints are per region or per resource.",
    "capabilities": [
      "speech.stt",
      "speech.streaming",
      "speech.batch",
      "speech.diarisation",
      "speech.languages",
      "speech.translation"
    ],
    "connect": {
      "http": "curl -X POST \"https://$AZURE_SPEECH_RESOURCE.cognitiveservices.azure.com/speechtotext/transcriptions:transcribe?api-version=2025-10-15\" \\\n  -H \"Ocp-Apim-Subscription-Key: $AZURE_SPEECH_KEY\" \\\n  -F \"audio=@call.wav\" -F 'definition={\"locales\":[\"en-US\"]}'",
      "install": "pip install azure-cognitiveservices-speech   # or: npm i microsoft-cognitiveservices-speech-sdk"
    },
    "docs": "https://learn.microsoft.com/en-us/azure/ai-services/speech-service/speech-to-text",
    "endpoint": "https://eastus.api.cognitive.microsoft.com/speechtotext",
    "grade": "BB",
    "json": "https://www.anchorterminal.com/api/v1/tools/azure-speech-to-text.json",
    "kind": "model",
    "leadPriceUsd": 0.00167,
    "leadUnit": "audio-minute",
    "letme": "https://letme.dev/azure-speech-to-text",
    "markdown": "https://www.anchorterminal.com/tools/azure-speech-to-text.md",
    "name": "Azure AI Speech speech-to-text",
    "priceSummary": "Freemium",
    "pricing": "freemium",
    "rank": 23,
    "sample": false,
    "score": 77,
    "slug": "azure-speech-to-text",
    "url": "https://www.anchorterminal.com/tools/azure-speech-to-text",
    "vendor": "Microsoft Azure",
    "where": "hosted",
    "x402": "no"
  },
  "picking": true,
  "rule": "highest Anchor grade for the capability, then price, compared only between prices whose item names the job, in one unit (the unit most listings at that grade price the job in, usage before plans), then p95 latency, then the Anchor score, then the name",
  "why": "Grade BB, the highest for speech.stt among listings that fit; price didn't decide (the rule compares only prices whose item names the job, and fewer than two listings at grade BB have one in the same unit), and it scores higher than Amazon Transcribe (77 against 73.6)."
}
