GOESB

librispeech-en-whisper-large-v3-batch

v1.0.0

openCC-BY-4.0scores against whisper-large-v3-en-batch

Audio

15 utterances · 145.195s total · 16000 Hz

Integrity

sha256: 088732d2cad9687abd258b16271d4fe5cc19dbb7370498d362f7f237f7984377

Full manifest

{
  "id": "librispeech-en-whisper-large-v3-batch",
  "version": "1.0.0",
  "sha256": "088732d2cad9687abd258b16271d4fe5cc19dbb7370498d362f7f237f7984377",
  "profile_id": "whisper-large-v3-en-batch",
  "visibility": "open",
  "license": "CC-BY-4.0",
  "audio": {
    "count": 15,
    "total_duration_s": 145.195,
    "sample_rate_hz": 16000,
    "manifest_sha256": "6b9cdcf81a565496b7d4ef4881b035d98593eb941e61d87780591f0aba5852d9",
    "source": {
      "type": "librispeech",
      "params": {
        "speaker": "1272",
        "chapter": "128104",
        "split": "dev-clean"
      },
      "fetch_instructions": "Identical audio to librispeech-en-batch — auto-fetched the same way. Manual fallback: run `python scripts/fetch_librispeech_subset.py --speaker 1272 --chapter 128104`, then pass `--audio-dir packs/librispeech-en-batch/audio` to `goesb run`."
    }
  },
  "metadata": {
    "language": "en-US",
    "recording_environment": "studio",
    "speech_style": "read",
    "transcription_style": "verbatim",
    "tags": [
      "librispeech",
      "en"
    ]
  }
}