pangkaiyu commited on
Commit
f97d188
·
verified ·
1 Parent(s): e059b82

Add files using upload-large-folder tool

Browse files
This view is limited to 50 files because it contains too many changes.   See raw diff
Files changed (50) hide show
  1. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1.jsonl +0 -0
  2. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1_default_performance.json +13 -0
  3. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1_wer_details.jsonl +0 -0
  4. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank0.log +5 -0
  5. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank1.log +2 -0
  6. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank2.log +2 -0
  7. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank3.log +2 -0
  8. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank4.log +2 -0
  9. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank5.log +2 -0
  10. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank6.log +2 -0
  11. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank7.log +2 -0
  12. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test.jsonl +0 -0
  13. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test_default_performance.json +137 -0
  14. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test_wer_details.jsonl +0 -0
  15. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank0.log +13 -0
  16. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank1.log +4 -0
  17. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank2.log +4 -0
  18. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank3.log +4 -0
  19. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank4.log +4 -0
  20. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank5.log +4 -0
  21. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank6.log +4 -0
  22. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank7.log +4 -0
  23. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/Qwen2.5-Omni-7B-lora2_voices_dev_test.jsonl +0 -0
  24. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/Qwen2.5-Omni-7B-lora2_voices_dev_test_wer_details.jsonl +0 -0
  25. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank1.log +4 -0
  26. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank2.log +4 -0
  27. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank3.log +4 -0
  28. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank4.log +4 -0
  29. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank5.log +4 -0
  30. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank6.log +4 -0
  31. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank7.log +4 -0
  32. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1.jsonl +0 -0
  33. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1_default_performance.json +13 -0
  34. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1_wer_details.jsonl +0 -0
  35. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank0.log +4 -0
  36. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank1.log +2 -0
  37. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank2.log +2 -0
  38. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank3.log +2 -0
  39. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank4.log +2 -0
  40. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank5.log +2 -0
  41. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank6.log +2 -0
  42. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank7.log +2 -0
  43. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test.jsonl +0 -0
  44. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test_default_performance.json +137 -0
  45. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test_wer_details.jsonl +0 -0
  46. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank0.log +9 -0
  47. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank1.log +4 -0
  48. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank2.log +4 -0
  49. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank3.log +4 -0
  50. different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank4.log +4 -0
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1_default_performance.json ADDED
@@ -0,0 +1,13 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "tedlium_test1",
4
+ "model": "Qwen2.5-Omni-7B-lora1",
5
+ "date": "2025-12-20 15:55:18.523324",
6
+ "performance": {
7
+ "TEDLIUM-test": {
8
+ "wer": 4.73,
9
+ "total": 1155
10
+ }
11
+ },
12
+ "eval_method": "qwen2-audio-impl"
13
+ }
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank0.log ADDED
@@ -0,0 +1,5 @@
 
 
 
 
 
 
1
+ 2025-12-20 13:24:56 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:56 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-16.11-27.84-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
3
+ 2025-12-20 13:26:32 | INFO | waiting for other ranks to finish, time elapsed: 10s
4
+ 2025-12-20 13:26:32 | INFO | model Qwen2.5-Omni-7B-lora1, data tedlium_test1, all 8 result merged to ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/Qwen2.5-Omni-7B-lora1_tedlium_test1.jsonl.
5
+ 2025-12-20 13:26:32 | INFO | skip eval for tedlium_test1
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:33 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:33 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-27.84-38.316-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:31 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:31 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-38.316-47.64-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:25 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:25 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-47.64-57.27-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:32 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:32 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-57.27-63.88-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:50 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:50 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-63.88-71.56-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:25 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:25 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-71.56-79.704-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/tedlium_test1/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:24:12 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: tedlium_test1
2
+ 2025-12-20 13:24:12 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-79.704-87.33-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test_default_performance.json ADDED
@@ -0,0 +1,137 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "voices_dev_test",
4
+ "model": "Qwen2.5-Omni-7B-lora1",
5
+ "date": "2025-12-20 15:55:17.661879",
6
+ "performance": {
7
+ "rm1-babb-clo": {
8
+ "wer": 3.39,
9
+ "total": 200
10
+ },
11
+ "rm1-babb-far": {
12
+ "wer": 6.33,
13
+ "total": 200
14
+ },
15
+ "rm1-musi-clo": {
16
+ "wer": 2.91,
17
+ "total": 200
18
+ },
19
+ "rm1-musi-far": {
20
+ "wer": 6.39,
21
+ "total": 200
22
+ },
23
+ "rm1-none-clo": {
24
+ "wer": 2.4,
25
+ "total": 200
26
+ },
27
+ "rm1-none-far": {
28
+ "wer": 3.07,
29
+ "total": 200
30
+ },
31
+ "rm1-tele-clo": {
32
+ "wer": 3.24,
33
+ "total": 200
34
+ },
35
+ "rm1-tele-far": {
36
+ "wer": 5.18,
37
+ "total": 200
38
+ },
39
+ "rm2-babb-clo": {
40
+ "wer": 3.18,
41
+ "total": 200
42
+ },
43
+ "rm2-babb-far": {
44
+ "wer": 5.98,
45
+ "total": 200
46
+ },
47
+ "rm2-musi-clo": {
48
+ "wer": 3.08,
49
+ "total": 200
50
+ },
51
+ "rm2-musi-far": {
52
+ "wer": 4.73,
53
+ "total": 200
54
+ },
55
+ "rm2-none-clo": {
56
+ "wer": 2.54,
57
+ "total": 200
58
+ },
59
+ "rm2-none-far": {
60
+ "wer": 3.21,
61
+ "total": 200
62
+ },
63
+ "rm2-tele-clo": {
64
+ "wer": 3.23,
65
+ "total": 200
66
+ },
67
+ "rm2-tele-far": {
68
+ "wer": 5.56,
69
+ "total": 200
70
+ },
71
+ "rm3-babb-clo": {
72
+ "wer": 14.62,
73
+ "total": 200
74
+ },
75
+ "rm3-babb-far": {
76
+ "wer": 108.57,
77
+ "total": 200
78
+ },
79
+ "rm3-musi-clo": {
80
+ "wer": 10.69,
81
+ "total": 200
82
+ },
83
+ "rm3-musi-far": {
84
+ "wer": 88.68,
85
+ "total": 200
86
+ },
87
+ "rm3-none-clo": {
88
+ "wer": 4.65,
89
+ "total": 200
90
+ },
91
+ "rm3-none-far": {
92
+ "wer": 23.81,
93
+ "total": 200
94
+ },
95
+ "rm3-tele-clo": {
96
+ "wer": 14.09,
97
+ "total": 200
98
+ },
99
+ "rm3-tele-far": {
100
+ "wer": 84.07,
101
+ "total": 200
102
+ },
103
+ "rm4-babb-clo": {
104
+ "wer": 3.91,
105
+ "total": 200
106
+ },
107
+ "rm4-babb-far": {
108
+ "wer": 109.37,
109
+ "total": 200
110
+ },
111
+ "rm4-musi-clo": {
112
+ "wer": 2.99,
113
+ "total": 200
114
+ },
115
+ "rm4-musi-far": {
116
+ "wer": 28.05,
117
+ "total": 200
118
+ },
119
+ "rm4-none-clo": {
120
+ "wer": 2.7,
121
+ "total": 200
122
+ },
123
+ "rm4-none-far": {
124
+ "wer": 4.65,
125
+ "total": 200
126
+ },
127
+ "rm4-tele-clo": {
128
+ "wer": 2.8,
129
+ "total": 200
130
+ },
131
+ "rm4-tele-far": {
132
+ "wer": 30.35,
133
+ "total": 200
134
+ }
135
+ },
136
+ "eval_method": "qwen2-audio-impl"
137
+ }
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank0.log ADDED
@@ -0,0 +1,13 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ 2025-12-20 13:11:10 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:10 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032639-sg0029-mc01-stu-clo-dg160.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:11:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
5
+ 2025-12-20 13:23:56 | INFO | waiting for other ranks to finish, time elapsed: 10s
6
+ 2025-12-20 13:24:06 | INFO | waiting for other ranks to finish, time elapsed: 20s
7
+ 2025-12-20 13:24:16 | INFO | waiting for other ranks to finish, time elapsed: 30s
8
+ 2025-12-20 13:24:26 | INFO | waiting for other ranks to finish, time elapsed: 40s
9
+ 2025-12-20 13:24:36 | INFO | waiting for other ranks to finish, time elapsed: 50s
10
+ 2025-12-20 13:24:46 | INFO | waiting for other ranks to finish, time elapsed: 60s
11
+ 2025-12-20 13:24:56 | INFO | waiting for other ranks to finish, time elapsed: 70s
12
+ 2025-12-20 13:24:56 | INFO | model Qwen2.5-Omni-7B-lora1, data voices_dev_test, all 8 result merged to ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/Qwen2.5-Omni-7B-lora1_voices_dev_test.jsonl.
13
+ 2025-12-20 13:24:56 | INFO | skip eval for voices_dev_test
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank1.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:12 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:12 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032639-sg0029-mc05-stu-far-dg160.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:11:12 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank2.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:10 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:10 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc05-stu-far-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:11:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank3.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:08 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:08 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc01-stu-clo-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:11:08 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank4.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:12 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:12 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130550-sg0026-mc01-stu-clo-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:11:13 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank5.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:10 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:10 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130550-sg0026-mc05-stu-far-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:11:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank6.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:10 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:10 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130551-sg0027-mc01-stu-clo-dg140.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:11:10 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora1/voices_dev_test/logs/rank7.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:11:12 | INFO | Running Qwen2.5-Omni-7B-lora1 on dataset: voices_dev_test
2
+ 2025-12-20 13:11:12 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130551-sg0027-mc05-stu-far-dg140.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:11:12 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/Qwen2.5-Omni-7B-lora2_voices_dev_test.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/Qwen2.5-Omni-7B-lora2_voices_dev_test_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank1.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:15 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:15 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032639-sg0029-mc05-stu-far-dg160.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:27:15 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank2.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:11 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:11 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc05-stu-far-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:27:12 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank3.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:15 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:15 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc01-stu-clo-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:27:16 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank4.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:15 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:15 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130550-sg0026-mc01-stu-clo-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:27:15 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank5.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:14 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:14 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130550-sg0026-mc05-stu-far-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:27:14 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank6.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:15 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:15 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130551-sg0027-mc01-stu-clo-dg140.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:27:15 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora2/voices_dev_test/logs/rank7.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:27:15 | INFO | Running Qwen2.5-Omni-7B-lora2 on dataset: voices_dev_test
2
+ 2025-12-20 13:27:15 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130551-sg0027-mc05-stu-far-dg140.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:27:15 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1_default_performance.json ADDED
@@ -0,0 +1,13 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "tedlium_test1",
4
+ "model": "Qwen2.5-Omni-7B-lora3",
5
+ "date": "2025-12-20 15:56:06.439401",
6
+ "performance": {
7
+ "TEDLIUM-test": {
8
+ "wer": 4.34,
9
+ "total": 1155
10
+ }
11
+ },
12
+ "eval_method": "qwen2-audio-impl"
13
+ }
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank0.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:57:05 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:57:05 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-16.11-27.84-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
3
+ 2025-12-20 13:58:43 | INFO | model Qwen2.5-Omni-7B-lora3, data tedlium_test1, all 8 result merged to ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/Qwen2.5-Omni-7B-lora3_tedlium_test1.jsonl.
4
+ 2025-12-20 13:58:43 | INFO | skip eval for tedlium_test1
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank1.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:39 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:39 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-27.84-38.316-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank2.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:35 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:35 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-38.316-47.64-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank3.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:48 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:48 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-47.64-57.27-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank4.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:36 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:36 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-57.27-63.88-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank5.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:48 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:48 | INFO | Msg example: {'index': 5, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-63.88-71.56-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank6.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:57:05 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:57:05 | INFO | Msg example: {'index': 6, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-71.56-79.704-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/tedlium_test1/logs/rank7.log ADDED
@@ -0,0 +1,2 @@
 
 
 
1
+ 2025-12-20 13:56:53 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: tedlium_test1
2
+ 2025-12-20 13:56:53 | INFO | Msg example: {'index': 7, 'audio': ['/workspace/intern/pangkaiyu/dg/tedlium_release1_data/test/wav/MichaelSpecter-79.704-87.33-<o,f0,male>.wav'], 'text': 'Please transcribe the audio content into text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'tedlium', 'dataset_name': 'tedlium_test1', 'lang': 'en', 'subset': 'TEDLIUM-test'}}
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test_default_performance.json ADDED
@@ -0,0 +1,137 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "task": "ASR",
3
+ "dataset": "voices_dev_test",
4
+ "model": "Qwen2.5-Omni-7B-lora3",
5
+ "date": "2025-12-20 15:56:05.566245",
6
+ "performance": {
7
+ "rm1-babb-clo": {
8
+ "wer": 3.52,
9
+ "total": 200
10
+ },
11
+ "rm1-babb-far": {
12
+ "wer": 6.2,
13
+ "total": 200
14
+ },
15
+ "rm1-musi-clo": {
16
+ "wer": 3.38,
17
+ "total": 200
18
+ },
19
+ "rm1-musi-far": {
20
+ "wer": 4.59,
21
+ "total": 200
22
+ },
23
+ "rm1-none-clo": {
24
+ "wer": 2.89,
25
+ "total": 200
26
+ },
27
+ "rm1-none-far": {
28
+ "wer": 3.46,
29
+ "total": 200
30
+ },
31
+ "rm1-tele-clo": {
32
+ "wer": 3.53,
33
+ "total": 200
34
+ },
35
+ "rm1-tele-far": {
36
+ "wer": 5.48,
37
+ "total": 200
38
+ },
39
+ "rm2-babb-clo": {
40
+ "wer": 3.7,
41
+ "total": 200
42
+ },
43
+ "rm2-babb-far": {
44
+ "wer": 6.68,
45
+ "total": 200
46
+ },
47
+ "rm2-musi-clo": {
48
+ "wer": 3.46,
49
+ "total": 200
50
+ },
51
+ "rm2-musi-far": {
52
+ "wer": 4.51,
53
+ "total": 200
54
+ },
55
+ "rm2-none-clo": {
56
+ "wer": 3.06,
57
+ "total": 200
58
+ },
59
+ "rm2-none-far": {
60
+ "wer": 3.42,
61
+ "total": 200
62
+ },
63
+ "rm2-tele-clo": {
64
+ "wer": 3.56,
65
+ "total": 200
66
+ },
67
+ "rm2-tele-far": {
68
+ "wer": 5.73,
69
+ "total": 200
70
+ },
71
+ "rm3-babb-clo": {
72
+ "wer": 14.07,
73
+ "total": 200
74
+ },
75
+ "rm3-babb-far": {
76
+ "wer": 85.99,
77
+ "total": 200
78
+ },
79
+ "rm3-musi-clo": {
80
+ "wer": 10.57,
81
+ "total": 200
82
+ },
83
+ "rm3-musi-far": {
84
+ "wer": 64.8,
85
+ "total": 200
86
+ },
87
+ "rm3-none-clo": {
88
+ "wer": 4.9,
89
+ "total": 200
90
+ },
91
+ "rm3-none-far": {
92
+ "wer": 23.18,
93
+ "total": 200
94
+ },
95
+ "rm3-tele-clo": {
96
+ "wer": 10.54,
97
+ "total": 200
98
+ },
99
+ "rm3-tele-far": {
100
+ "wer": 77.83,
101
+ "total": 200
102
+ },
103
+ "rm4-babb-clo": {
104
+ "wer": 4.16,
105
+ "total": 200
106
+ },
107
+ "rm4-babb-far": {
108
+ "wer": 103.69,
109
+ "total": 200
110
+ },
111
+ "rm4-musi-clo": {
112
+ "wer": 3.12,
113
+ "total": 200
114
+ },
115
+ "rm4-musi-far": {
116
+ "wer": 23.01,
117
+ "total": 200
118
+ },
119
+ "rm4-none-clo": {
120
+ "wer": 2.73,
121
+ "total": 200
122
+ },
123
+ "rm4-none-far": {
124
+ "wer": 4.92,
125
+ "total": 200
126
+ },
127
+ "rm4-tele-clo": {
128
+ "wer": 3.08,
129
+ "total": 200
130
+ },
131
+ "rm4-tele-far": {
132
+ "wer": 25.78,
133
+ "total": 200
134
+ }
135
+ },
136
+ "eval_method": "qwen2-audio-impl"
137
+ }
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test_wer_details.jsonl ADDED
The diff for this file is too large to render. See raw diff
 
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank0.log ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ 2025-12-20 13:43:24 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: voices_dev_test
2
+ 2025-12-20 13:43:24 | INFO | Msg example: {'index': 0, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032639-sg0029-mc01-stu-clo-dg160.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:43:25 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
5
+ 2025-12-20 13:56:45 | INFO | waiting for other ranks to finish, time elapsed: 10s
6
+ 2025-12-20 13:56:55 | INFO | waiting for other ranks to finish, time elapsed: 20s
7
+ 2025-12-20 13:57:05 | INFO | waiting for other ranks to finish, time elapsed: 30s
8
+ 2025-12-20 13:57:05 | INFO | model Qwen2.5-Omni-7B-lora3, data voices_dev_test, all 8 result merged to ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/Qwen2.5-Omni-7B-lora3_voices_dev_test.jsonl.
9
+ 2025-12-20 13:57:05 | INFO | skip eval for voices_dev_test
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank1.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:43:24 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: voices_dev_test
2
+ 2025-12-20 13:43:24 | INFO | Msg example: {'index': 1, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032639-sg0029-mc05-stu-far-dg160.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:43:24 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank2.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:43:23 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: voices_dev_test
2
+ 2025-12-20 13:43:23 | INFO | Msg example: {'index': 2, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc05-stu-far-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-far'}}
3
+ 2025-12-20 13:43:24 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank3.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:43:23 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: voices_dev_test
2
+ 2025-12-20 13:43:23 | INFO | Msg example: {'index': 3, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp4899/Lab41-SRI-VOiCES-rm4-babb-sp4899-ch032658-sg0012-mc01-stu-clo-dg070.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:43:24 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.
different_distribution/ff_5e-6_new_linear/Qwen2.5-Omni-7B-lora3/voices_dev_test/logs/rank4.log ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ 2025-12-20 13:43:24 | INFO | Running Qwen2.5-Omni-7B-lora3 on dataset: voices_dev_test
2
+ 2025-12-20 13:43:24 | INFO | Msg example: {'index': 4, 'audio': ['/workspace/intern/pangkaiyu/dg/VOiCES_devkit/distant-16k/speech/test/rm4/babb/sp1447/Lab41-SRI-VOiCES-rm4-babb-sp1447-ch130550-sg0026-mc01-stu-clo-dg010.wav'], 'text': 'Please transcribe the spoken content into written text.', 'meta': {'task': 'ASR', 'interactive': 'Audio-analysis', 'audio_type': 'Speech', 'dataset_series': 'voices_dev', 'dataset_name': 'voices_dev_test', 'lang': 'en', 'subset': 'rm4-babb-clo'}}
3
+ 2025-12-20 13:43:24 | INFO | Prompt: You are a speech recognition model.
4
+ Transcribe the English audio into text without any punctuation marks.