{"schemaVersion":"1.0","generatedFrom":"https://brightaifuture.com/discoveries/whisper-local-transcription","record":{"id":"whisper-local-transcription","headline":"Transcribing speech on infrastructure you control","canonicalUrl":"https://brightaifuture.com/discoveries/whisper-local-transcription","datePublished":"2026-09-19","dateModified":null,"sourcePublicationDate":"2022-09-21","author":null,"publisher":{"name":"Bright AI Future","url":"https://brightaifuture.com/"},"topics":["open-models"],"summary":"Whisper provides downloadable speech-recognition models for transcription, language identification, translation into English, and caption-making without requiring a hosted speech service.","evidenceState":"Demonstrated","keyFacts":[{"label":"AI’s role","value":"A multilingual sequence-to-sequence model converts audio into timestamped text or translated text."},{"label":"Documented result","value":"OpenAI released model weights, inference code, and a command-line interface under MIT terms in September 2022, making local runs broadly reproducible."},{"label":"Important limitation","value":"Accuracy varies with language, accent, noise, recording conditions, and subject matter. The training corpus and a complete training recipe were not released."}],"limitations":["Accuracy varies with language, accent, noise, recording conditions, and subject matter. The training corpus and a complete training recipe were not released.","The official repository establishes MIT-licensed code and weights. That unusually permissive artifact release remains distinct from access to the training data."],"evidenceLinks":[{"title":"Whisper","url":"https://github.com/openai/whisper","type":"repository"},{"title":"Introducing Whisper","url":"https://openai.com/index/whisper/","type":"institution"}],"evidencePackUrl":"https://brightaifuture.com/evidence-pack/whisper-local-transcription","embedUrl":"https://brightaifuture.com/embed/story/whisper-local-transcription","attribution":{"credit":"Bright AI Future","requirements":["Link to the canonical Bright record.","Keep material limitations with the claim they qualify.","Link to the original evidence when repeating a substantive claim.","Do not describe a source check or organization-reported result as independent verification."],"sourceRights":"Linked source material, quotations, trademarks and media remain subject to their owners’ terms. No reuse right is granted for third-party media."}},"claim":{"humanProblem":"People need searchable transcripts and captions, including when audio is private or connectivity is limited.","priorConstraint":"Strong speech recognition often depended on a remote service or a model tuned to a narrow language and acoustic setting.","aiRole":"A multilingual sequence-to-sequence model converts audio into timestamped text or translated text.","documentedResult":"OpenAI released model weights, inference code, and a command-line interface under MIT terms in September 2022, making local runs broadly reproducible.","whyItMayMatter":"Local weights give organizations more control over where audio travels, but people still need ways to inspect names, meaning, and accessibility quality.","unresolvedQuestions":[]},"evidenceAssessment":{"state":"Demonstrated","claimConfidence":"unassessed","reviewState":"source-checked","reviewMethod":"ai-assisted","reviewNote":"AI-assisted comparison with the cited sources. Source-checked means the record was checked against those sources; it does not claim independent reproduction, expert review, or validation of the publisher’s results.","lastSourceReview":"2026-09-19","independentVerification":"not-established-by-this-source-review"},"sources":[{"id":"whisper-repository","title":"Whisper","url":"https://github.com/openai/whisper","type":"repository"},{"id":"whisper-release","title":"Introducing Whisper","url":"https://openai.com/index/whisper/","type":"institution"}],"revisions":[{"id":"revision:open-models-added:whisper-local-transcription","recordedAt":"2026-09-19","summary":"Bright added this source-checked open-model application record. The cited source publication date is 2022-09-21; 2026-09-19 is when Bright added this record.","sourceIds":["whisper-repository","whisper-release"]}],"corrections":[]}