GOESB

librispeech-en-whisper-tiny-batch

v1.0.0

openCC-BY-4.0scores against whisper-tiny-en-batch

Audio

15 utterances · 145.195s total · 16000 Hz

Integrity

sha256: 59a1e72041f262e6a3ff2837fc14c0b725ef5d674fb6433f9e228a4b2e3c5bd2

Full manifest

{
  "id": "librispeech-en-whisper-tiny-batch",
  "version": "1.0.0",
  "sha256": "59a1e72041f262e6a3ff2837fc14c0b725ef5d674fb6433f9e228a4b2e3c5bd2",
  "profile_id": "whisper-tiny-en-batch",
  "visibility": "open",
  "license": "CC-BY-4.0",
  "audio": {
    "count": 15,
    "total_duration_s": 145.195,
    "sample_rate_hz": 16000,
    "manifest_sha256": "6b9cdcf81a565496b7d4ef4881b035d98593eb941e61d87780591f0aba5852d9",
    "source": {
      "type": "librispeech",
      "params": {
        "speaker": "1272",
        "chapter": "128104",
        "split": "dev-clean"
      },
      "fetch_instructions": "Identical audio to librispeech-en-batch — auto-fetched the same way. Manual fallback: run `python scripts/fetch_librispeech_subset.py --speaker 1272 --chapter 128104`, then pass `--audio-dir packs/librispeech-en-batch/audio` to `goesb run`."
    }
  },
  "metadata": {
    "language": "en-US",
    "recording_environment": "studio",
    "speech_style": "read",
    "transcription_style": "verbatim",
    "tags": [
      "librispeech",
      "en"
    ]
  }
}