librispeech-en-whispercpp
v1.0.0
openCC-BY-4.0en-US
Usable by any en-US profile — see results using this pack.
Audio
15 utterances · 145.195s total · 16000 Hz
Integrity
sha256: f6c58f6988e81bb847a226dc37af3b350ffbfa1846cd7700f8720df75620b30d
Full manifest
{
"id": "librispeech-en-whispercpp",
"version": "1.0.0",
"sha256": "f6c58f6988e81bb847a226dc37af3b350ffbfa1846cd7700f8720df75620b30d",
"profile_id": "whispercpp-base-en-batch",
"visibility": "open",
"license": "CC-BY-4.0",
"audio": {
"count": 15,
"total_duration_s": 145.195,
"sample_rate_hz": 16000,
"manifest_sha256": "6b9cdcf81a565496b7d4ef4881b035d98593eb941e61d87780591f0aba5852d9",
"source": {
"type": "librispeech",
"params": {
"speaker": "1272",
"chapter": "128104",
"split": "dev-clean"
},
"fetch_instructions": "Identical audio to librispeech-en — auto-fetched the same way (same speaker/chapter). Manual fallback: run scripts/fetch_librispeech_subset.py for that pack, then pass `--audio-dir packs/librispeech-en/audio` to `goesb run`."
}
},
"metadata": {
"language": "en-US",
"dialect": "general-american",
"age_group": "mixed",
"recording_environment": "quiet",
"microphone": "consumer usb",
"sample_rate_hz": 16000,
"background_noise": "none",
"num_speakers": 100,
"speech_style": "read",
"transcription_style": "verbatim",
"tags": [
"librispeech",
"english",
"clean",
"whisper-cpp"
]
}
}