File size: 2,747 Bytes
ed957e6 | 1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30 31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48 49 50 51 52 53 54 55 56 57 58 59 60 61 62 63 64 65 66 67 68 69 70 71 72 73 74 75 76 77 78 79 80 81 82 83 84 85 86 | {
"passed": true,
"checked_at_utc": "2026-09-27T02:40:28.055408+00:00",
"scope": "DuDE training data; upstream Qwen pretraining coverage is not known",
"examples": [
{
"label": "Sample 1",
"conversation_id": "V01_S0192_I00000307",
"example_id": "V01_S0192_I00000307-0000",
"split": "dev",
"participants": {
"A": "V01_P1356",
"B": "V01_P1357"
},
"transcript_source": "own_Qwen3-ASR+ForcedAligner+PANNs",
"dataset_transcript_downloaded": false,
"annotation_sha256": "76474198f55a6d7afd7f826cd910b9841af8293a32413f2f8b672c1ef7fd4996",
"public_example_id": "sample_1"
},
{
"label": "Sample 2",
"conversation_id": "V01_S0220_I00000126",
"example_id": "V01_S0220_I00000126-0000",
"split": "dev",
"participants": {
"A": "V01_P1409",
"B": "V01_P1500"
},
"transcript_source": "own_Qwen3-ASR+ForcedAligner+PANNs",
"dataset_transcript_downloaded": false,
"annotation_sha256": "d3d434ee2d09aa3916f1883df7e1b1134125b285a2a09cac64171c49110d1cb6",
"public_example_id": "sample_2"
}
],
"training_manifests": [
{
"corpus": "seamless-100h",
"train_conversations": 1468,
"sha256": "05c7e77c69078b9d68ef3c30a68392d893f82f7a5cca94bdbf2e3a4ced40d282",
"conversation_matches": [],
"participant_matches": [],
"featured_conversations_in_dev": [
"V01_S0220_I00000126"
]
},
{
"corpus": "seamless-500h-v1",
"train_conversations": 7365,
"sha256": "2967b0c49574182fd14d56571107dbc264f199d6cd3ed39ea71bc9956a81203a",
"conversation_matches": [],
"participant_matches": [],
"featured_conversations_in_dev": [
"V01_S0192_I00000307",
"V01_S0220_I00000126"
]
},
{
"corpus": "seamless-2000h-v1",
"train_conversations": 31493,
"sha256": "08d437e07c00ebd4a73feed219e0048e69d1bdcb681dac12f20b0fd9c803758e",
"conversation_matches": [],
"participant_matches": [],
"featured_conversations_in_dev": [
"V01_S0192_I00000307",
"V01_S0220_I00000126"
]
}
],
"training_codec_caches": [
{
"corpus": "conversation-2000h-v1",
"shards": 8,
"examples": 30854,
"featured_conversation_matches": []
},
{
"corpus": "conversation-500h-full-v1",
"shards": 4,
"examples": 7190,
"featured_conversation_matches": []
}
],
"transcript_note": "XML uses our ASR/event annotations of held-out Seamless Interaction audio, not dataset-provided transcripts.",
"voice_reference_note": "Speaker reference clips are provided during inference; they are not training examples."
}
|