pangkaiyu commited on
Commit
0d237b2
·
verified ·
1 Parent(s): 7739acc

Upload folder using huggingface_hub

Browse files
Files changed (33) hide show
  1. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl +0 -0
  2. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_default_performance.json +17 -0
  3. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_wer_details.jsonl +0 -0
  4. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank0.log +9 -0
  5. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank1.log +4 -0
  6. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank2.log +4 -0
  7. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank3.log +4 -0
  8. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank4.log +4 -0
  9. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank5.log +4 -0
  10. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank6.log +4 -0
  11. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank7.log +4 -0
  12. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl +0 -0
  13. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_default_performance.json +121 -0
  14. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_wer_details.jsonl +0 -0
  15. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank0.log +4 -0
  16. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank1.log +2 -0
  17. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank2.log +2 -0
  18. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank3.log +2 -0
  19. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank4.log +2 -0
  20. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank5.log +2 -0
  21. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank6.log +2 -0
  22. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank7.log +2 -0
  23. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl +0 -0
  24. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_default_performance.json +25 -0
  25. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_wer_details.jsonl +0 -0
  26. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank0.log +5 -0
  27. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank1.log +2 -0
  28. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank2.log +2 -0
  29. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank3.log +2 -0
  30. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank4.log +2 -0
  31. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank5.log +2 -0
  32. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank6.log +2 -0
  33. 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank7.log +2 -0
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_default_performance.json ADDED
@@ -0,0 +1,17 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "LibriSpeech",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 07:37:32.640411",
6
+ "performance": {
7
+ "test_clean": {
8
+ "wer": 2.61,
9
+ "total": 2620
10
+ },
11
+ "test_other": {
12
+ "wer": 4.7,
13
+ "total": 2939
14
+ }
15
+ },
16
+ "eval_method": "qwen2-audio-impl"
17
+ }
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank0.log ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ 2025-12-17 07:24:07 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:07 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0003.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:08 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
5
+ 2025-12-17 07:31:01 | INFO | waiting for other ranks to finish, time elapsed: 10s
6
+ 2025-12-17 07:31:11 | INFO | waiting for other ranks to finish, time elapsed: 20s
7
+ 2025-12-17 07:31:21 | INFO | waiting for other ranks to finish, time elapsed: 30s
8
+ 2025-12-17 07:31:22 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data LibriSpeech, all 8 result merged to 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl.
9
+ 2025-12-17 07:31:22 | INFO | skip eval for LibriSpeech
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank1.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:11 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:11 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0012.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:12 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank2.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:09 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:09 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0026.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank3.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:09 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:09 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0022.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank4.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:09 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:09 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0002.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank5.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:11 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:11 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0025.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:11 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank6.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:10 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:10 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0010.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:11 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank7.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:24:10 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 07:24:10 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0001.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 07:24:11 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_default_performance.json ADDED
@@ -0,0 +1,121 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "noizeus",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 07:37:33.016730",
6
+ "performance": {
7
+ "airport_0dB": {
8
+ "wer": 27.27,
9
+ "total": 30
10
+ },
11
+ "airport_10dB": {
12
+ "wer": 1.65,
13
+ "total": 30
14
+ },
15
+ "airport_15dB": {
16
+ "wer": 2.07,
17
+ "total": 30
18
+ },
19
+ "airport_5dB": {
20
+ "wer": 11.98,
21
+ "total": 30
22
+ },
23
+ "babble_0dB": {
24
+ "wer": 42.56,
25
+ "total": 30
26
+ },
27
+ "babble_10dB": {
28
+ "wer": 1.65,
29
+ "total": 30
30
+ },
31
+ "babble_15dB": {
32
+ "wer": 2.48,
33
+ "total": 30
34
+ },
35
+ "babble_5dB": {
36
+ "wer": 11.16,
37
+ "total": 30
38
+ },
39
+ "car_0dB": {
40
+ "wer": 40.91,
41
+ "total": 30
42
+ },
43
+ "car_10dB": {
44
+ "wer": 4.96,
45
+ "total": 30
46
+ },
47
+ "car_15dB": {
48
+ "wer": 1.24,
49
+ "total": 30
50
+ },
51
+ "car_5dB": {
52
+ "wer": 11.57,
53
+ "total": 30
54
+ },
55
+ "exhibition_0dB": {
56
+ "wer": 28.1,
57
+ "total": 30
58
+ },
59
+ "exhibition_10dB": {
60
+ "wer": 2.89,
61
+ "total": 30
62
+ },
63
+ "exhibition_15dB": {
64
+ "wer": 3.31,
65
+ "total": 30
66
+ },
67
+ "exhibition_5dB": {
68
+ "wer": 9.09,
69
+ "total": 30
70
+ },
71
+ "restaurant_0dB": {
72
+ "wer": 37.6,
73
+ "total": 30
74
+ },
75
+ "restaurant_10dB": {
76
+ "wer": 3.72,
77
+ "total": 30
78
+ },
79
+ "restaurant_15dB": {
80
+ "wer": 1.65,
81
+ "total": 30
82
+ },
83
+ "restaurant_5dB": {
84
+ "wer": 13.22,
85
+ "total": 30
86
+ },
87
+ "station_0dB": {
88
+ "wer": 35.12,
89
+ "total": 30
90
+ },
91
+ "station_10dB": {
92
+ "wer": 2.48,
93
+ "total": 30
94
+ },
95
+ "station_15dB": {
96
+ "wer": 0.83,
97
+ "total": 30
98
+ },
99
+ "station_5dB": {
100
+ "wer": 9.92,
101
+ "total": 30
102
+ },
103
+ "street_0dB": {
104
+ "wer": 36.36,
105
+ "total": 30
106
+ },
107
+ "street_10dB": {
108
+ "wer": 4.55,
109
+ "total": 30
110
+ },
111
+ "street_15dB": {
112
+ "wer": 2.07,
113
+ "total": 30
114
+ },
115
+ "street_5dB": {
116
+ "wer": 15.7,
117
+ "total": 30
118
+ }
119
+ },
120
+ "eval_method": "qwen2-audio-impl"
121
+ }
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank0.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 07:31:22 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:22 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp01_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
3
+ 2025-12-17 07:31:59 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data noizeus, all 8 result merged to 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl.
4
+ 2025-12-17 07:31:59 | INFO | skip eval for noizeus
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:01 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:01 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp02_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:30:58 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:30:58 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp03_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:16 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:16 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp04_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:01 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:01 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp05_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:05 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:05 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp06_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:30:44 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:30:44 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp07_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:13 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 07:31:13 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp08_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_default_performance.json ADDED
@@ -0,0 +1,25 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "voices_dev_clo",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 07:37:35.205757",
6
+ "performance": {
7
+ "babb": {
8
+ "wer": 3.77,
9
+ "total": 364
10
+ },
11
+ "musi": {
12
+ "wer": 3.14,
13
+ "total": 371
14
+ },
15
+ "none": {
16
+ "wer": 2.53,
17
+ "total": 380
18
+ },
19
+ "tele": {
20
+ "wer": 3.27,
21
+ "total": 351
22
+ }
23
+ },
24
+ "eval_method": "qwen2-audio-impl"
25
+ }
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank0.log ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ 2025-12-17 07:31:59 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:59 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0112/Lab41-SRI-VOiCES-rm1-babb-sp0112-ch123215-sg0025-mc01-stu-clo-dg080.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
3
+ 2025-12-17 07:35:27 | INFO | waiting for other ranks to finish, time elapsed: 10s
4
+ 2025-12-17 07:35:27 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data voices_dev_clo, all 8 result merged to 7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl.
5
+ 2025-12-17 07:35:27 | INFO | skip eval for voices_dev_clo
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:39 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:39 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0122/Lab41-SRI-VOiCES-rm1-babb-sp0122-ch121729-sg0002-mc02-lav-clo-dg060.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:37 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:37 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0122/Lab41-SRI-VOiCES-rm1-babb-sp0122-ch121730-sg0014-mc01-stu-clo-dg000.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:54 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:54 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0159/Lab41-SRI-VOiCES-rm1-babb-sp0159-ch135897-sg0052-mc01-stu-clo-dg100.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:40 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:40 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0174/Lab41-SRI-VOiCES-rm1-babb-sp0174-ch084280-sg0013-mc02-lav-clo-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:43 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:43 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0188/Lab41-SRI-VOiCES-rm1-babb-sp0188-ch135249-sg0029-mc01-stu-clo-dg170.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:23 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:23 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0205/Lab41-SRI-VOiCES-rm1-babb-sp0205-ch159056-sg0032-mc01-stu-clo-dg020.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck300/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 07:31:52 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 07:31:52 | INFO | Msg example: {'index': 8, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0208/Lab41-SRI-VOiCES-rm1-babb-sp0208-ch126851-sg0011-mc02-lav-clo-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}