pangkaiyu commited on
Commit
7739acc
·
verified ·
1 Parent(s): 14d3c20

Upload folder using huggingface_hub

Browse files
Files changed (33) hide show
  1. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl +0 -0
  2. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_default_performance.json +17 -0
  3. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_wer_details.jsonl +0 -0
  4. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank0.log +10 -0
  5. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank1.log +4 -0
  6. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank2.log +4 -0
  7. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank3.log +4 -0
  8. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank4.log +4 -0
  9. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank5.log +4 -0
  10. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank6.log +4 -0
  11. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank7.log +4 -0
  12. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl +0 -0
  13. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_default_performance.json +121 -0
  14. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_wer_details.jsonl +0 -0
  15. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank0.log +4 -0
  16. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank1.log +2 -0
  17. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank2.log +2 -0
  18. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank3.log +2 -0
  19. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank4.log +2 -0
  20. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank5.log +2 -0
  21. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank6.log +2 -0
  22. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank7.log +2 -0
  23. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl +0 -0
  24. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_default_performance.json +25 -0
  25. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_wer_details.jsonl +0 -0
  26. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank0.log +5 -0
  27. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank1.log +2 -0
  28. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank2.log +2 -0
  29. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank3.log +2 -0
  30. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank4.log +2 -0
  31. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank5.log +2 -0
  32. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank6.log +2 -0
  33. 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank7.log +2 -0
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_default_performance.json ADDED
@@ -0,0 +1,17 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "LibriSpeech",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 09:39:59.362503",
6
+ "performance": {
7
+ "test_clean": {
8
+ "wer": 2.53,
9
+ "total": 2620
10
+ },
11
+ "test_other": {
12
+ "wer": 4.58,
13
+ "total": 2939
14
+ }
15
+ },
16
+ "eval_method": "qwen2-audio-impl"
17
+ }
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank0.log ADDED
@@ -0,0 +1,10 @@
 
 
 
 
 
 
 
 
 
 
 
1
+ 2025-12-17 08:42:17 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:17 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0003.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:18 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
5
+ 2025-12-17 08:48:45 | INFO | waiting for other ranks to finish, time elapsed: 10s
6
+ 2025-12-17 08:48:55 | INFO | waiting for other ranks to finish, time elapsed: 20s
7
+ 2025-12-17 08:49:05 | INFO | waiting for other ranks to finish, time elapsed: 30s
8
+ 2025-12-17 08:49:15 | INFO | waiting for other ranks to finish, time elapsed: 40s
9
+ 2025-12-17 08:49:15 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data LibriSpeech, all 8 result merged to 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/Qwen2.5-Omni-7B-l-after-ea_LibriSpeech.jsonl.
10
+ 2025-12-17 08:49:15 | INFO | skip eval for LibriSpeech
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank1.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:21 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:21 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0012.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:21 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank2.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:21 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:21 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0026.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:21 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank3.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:17 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:17 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0022.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:17 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank4.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:19 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:19 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0002.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:20 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank5.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:21 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:21 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0025.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:22 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank6.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:20 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:20 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0010.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:20 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/LibriSpeech/logs/rank7.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:42:20 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: LibriSpeech
2
+ 2025-12-17 08:42:20 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/dataset/LibriSpeech/librispeech/LibriSpeech/test-clean/1320/122617/1320-122617-0001.flac'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'LibriSpeech', 'dataset_name': 'LibriSpeech', 'lang': 'en', 'subset': 'test_clean'}}
3
+ 2025-12-17 08:42:21 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_default_performance.json ADDED
@@ -0,0 +1,121 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "noizeus",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 09:39:59.732647",
6
+ "performance": {
7
+ "airport_0dB": {
8
+ "wer": 26.86,
9
+ "total": 30
10
+ },
11
+ "airport_10dB": {
12
+ "wer": 1.65,
13
+ "total": 30
14
+ },
15
+ "airport_15dB": {
16
+ "wer": 1.65,
17
+ "total": 30
18
+ },
19
+ "airport_5dB": {
20
+ "wer": 11.57,
21
+ "total": 30
22
+ },
23
+ "babble_0dB": {
24
+ "wer": 42.98,
25
+ "total": 30
26
+ },
27
+ "babble_10dB": {
28
+ "wer": 3.31,
29
+ "total": 30
30
+ },
31
+ "babble_15dB": {
32
+ "wer": 2.48,
33
+ "total": 30
34
+ },
35
+ "babble_5dB": {
36
+ "wer": 9.5,
37
+ "total": 30
38
+ },
39
+ "car_0dB": {
40
+ "wer": 42.56,
41
+ "total": 30
42
+ },
43
+ "car_10dB": {
44
+ "wer": 4.13,
45
+ "total": 30
46
+ },
47
+ "car_15dB": {
48
+ "wer": 1.24,
49
+ "total": 30
50
+ },
51
+ "car_5dB": {
52
+ "wer": 12.4,
53
+ "total": 30
54
+ },
55
+ "exhibition_0dB": {
56
+ "wer": 30.17,
57
+ "total": 30
58
+ },
59
+ "exhibition_10dB": {
60
+ "wer": 3.31,
61
+ "total": 30
62
+ },
63
+ "exhibition_15dB": {
64
+ "wer": 2.89,
65
+ "total": 30
66
+ },
67
+ "exhibition_5dB": {
68
+ "wer": 10.74,
69
+ "total": 30
70
+ },
71
+ "restaurant_0dB": {
72
+ "wer": 39.26,
73
+ "total": 30
74
+ },
75
+ "restaurant_10dB": {
76
+ "wer": 2.48,
77
+ "total": 30
78
+ },
79
+ "restaurant_15dB": {
80
+ "wer": 2.07,
81
+ "total": 30
82
+ },
83
+ "restaurant_5dB": {
84
+ "wer": 11.16,
85
+ "total": 30
86
+ },
87
+ "station_0dB": {
88
+ "wer": 39.26,
89
+ "total": 30
90
+ },
91
+ "station_10dB": {
92
+ "wer": 2.89,
93
+ "total": 30
94
+ },
95
+ "station_15dB": {
96
+ "wer": 0.41,
97
+ "total": 30
98
+ },
99
+ "station_5dB": {
100
+ "wer": 10.33,
101
+ "total": 30
102
+ },
103
+ "street_0dB": {
104
+ "wer": 41.32,
105
+ "total": 30
106
+ },
107
+ "street_10dB": {
108
+ "wer": 4.96,
109
+ "total": 30
110
+ },
111
+ "street_15dB": {
112
+ "wer": 1.24,
113
+ "total": 30
114
+ },
115
+ "street_5dB": {
116
+ "wer": 15.7,
117
+ "total": 30
118
+ }
119
+ },
120
+ "eval_method": "qwen2-audio-impl"
121
+ }
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank0.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-17 08:49:15 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:15 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp01_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
3
+ 2025-12-17 08:49:49 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data noizeus, all 8 result merged to 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/Qwen2.5-Omni-7B-l-after-ea_noizeus.jsonl.
4
+ 2025-12-17 08:49:49 | INFO | skip eval for noizeus
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:10 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:10 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp02_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:05 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:05 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp03_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:05 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:05 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp04_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:04 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:04 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp05_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:01 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:01 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp06_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:48:52 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:48:52 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp07_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/noizeus/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:02 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: noizeus
2
+ 2025-12-17 08:49:02 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/Kimi-Audio/Kimi-Audio-Evalkit/data/downloaded_datasets/noizeus/noizeus/airport/0dB/sp08_airport_sn0.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'noizeus', 'dataset_name': 'noizeus', 'lang': 'en', 'subset': 'airport_0dB'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_default_performance.json ADDED
@@ -0,0 +1,25 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "voices_dev_clo",
4
+ "model": "Qwen2.5-Omni-7B-l-after-ea",
5
+ "date": "2025-12-17 09:40:01.900408",
6
+ "performance": {
7
+ "babb": {
8
+ "wer": 3.77,
9
+ "total": 364
10
+ },
11
+ "musi": {
12
+ "wer": 3.16,
13
+ "total": 371
14
+ },
15
+ "none": {
16
+ "wer": 3.91,
17
+ "total": 380
18
+ },
19
+ "tele": {
20
+ "wer": 3.23,
21
+ "total": 351
22
+ }
23
+ },
24
+ "eval_method": "qwen2-audio-impl"
25
+ }
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank0.log ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ 2025-12-17 08:49:49 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:49 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0112/Lab41-SRI-VOiCES-rm1-babb-sp0112-ch123215-sg0025-mc01-stu-clo-dg080.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
3
+ 2025-12-17 08:53:15 | INFO | waiting for other ranks to finish, time elapsed: 10s
4
+ 2025-12-17 08:53:15 | INFO | model Qwen2.5-Omni-7B-l-after-ea, data voices_dev_clo, all 8 result merged to 7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/Qwen2.5-Omni-7B-l-after-ea_voices_dev_clo.jsonl.
5
+ 2025-12-17 08:53:15 | INFO | skip eval for voices_dev_clo
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:47 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:47 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0122/Lab41-SRI-VOiCES-rm1-babb-sp0122-ch121729-sg0002-mc02-lav-clo-dg060.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:42 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:42 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0122/Lab41-SRI-VOiCES-rm1-babb-sp0122-ch121730-sg0014-mc01-stu-clo-dg000.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:43 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:43 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0159/Lab41-SRI-VOiCES-rm1-babb-sp0159-ch135897-sg0052-mc01-stu-clo-dg100.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:42 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:42 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0174/Lab41-SRI-VOiCES-rm1-babb-sp0174-ch084280-sg0013-mc02-lav-clo-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:38 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:38 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0188/Lab41-SRI-VOiCES-rm1-babb-sp0188-ch135249-sg0029-mc01-stu-clo-dg170.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:31 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:31 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0205/Lab41-SRI-VOiCES-rm1-babb-sp0205-ch159056-sg0032-mc01-stu-clo-dg020.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}
7B_l_after_ea300_ck200/Qwen2.5-Omni-7B-l-after-ea/voices_dev_clo/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-17 08:49:39 | INFO | Running Qwen2.5-Omni-7B-l-after-ea on dataset: voices_dev_clo
2
+ 2025-12-17 08:49:39 | INFO | Msg example: {'index': 8, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_Box_unzip/Development_Data/Automatic_Speech_Recognition/ASR_dev.v2/rm1/babb/sp_0032-1182/sp0208/Lab41-SRI-VOiCES-rm1-babb-sp0208-ch126851-sg0011-mc02-lav-clo-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices', 'dataset_name': 'voices_dev_clo', 'lang': 'en', 'subset': 'babb'}}