Update README.md
Browse files
README.md
CHANGED
@@ -1,3 +1,69 @@
|
|
1 |
-
|
2 |
-
|
3 |
-
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
## ASR Leaderboard
|
2 |
+
<table border="0" cellspacing="0" cellpadding="6" style="border-collapse:collapse;">
|
3 |
+
<tr>
|
4 |
+
<th align="left" rowspan="2">Model</th>
|
5 |
+
<th align="center" rowspan="2">#Params (M)</th>
|
6 |
+
<th align="center" colspan="2">In-House</th>
|
7 |
+
<th align="center" colspan="5">Open-Source</th>
|
8 |
+
<th align="center" colspan="2">WSYue-eval</th>
|
9 |
+
</tr>
|
10 |
+
<tr>
|
11 |
+
<th align="center">Dialogue</th>
|
12 |
+
<th align="center">Reading</th>
|
13 |
+
<th align="center">yue</th>
|
14 |
+
<th align="center">HK</th>
|
15 |
+
<th align="center">MDCC</th>
|
16 |
+
<th align="center">Daily_Use</th>
|
17 |
+
<th align="center">Commands</th>
|
18 |
+
<th align="center">Short</th>
|
19 |
+
<th align="center">Long</th>
|
20 |
+
</tr>
|
21 |
+
|
22 |
+
<tr><td align="left" colspan="11"><b>w/o LLM</b></td></tr>
|
23 |
+
<tr>
|
24 |
+
<td align="left"><b>Conformer-Yue⭐</b></td><td align="center">130</td><td align="center"><b>16.57</b></td><td align="center">7.82</td><td align="center">7.72</td><td align="center">11.42</td><td align="center">5.73</td><td align="center">5.73</td><td align="center">8.97</td><td align="center"><ins>5.05</ins></td><td align="center">8.89</td>
|
25 |
+
</tr>
|
26 |
+
<tr>
|
27 |
+
<td align="left">Paraformer</td><td align="center">220</td><td align="center">83.22</td><td align="center">51.97</td><td align="center">70.16</td><td align="center">68.49</td><td align="center">47.67</td><td align="center">79.31</td><td align="center">69.32</td><td align="center">73.64</td><td align="center">89.00</td>
|
28 |
+
</tr>
|
29 |
+
<tr>
|
30 |
+
<td align="left">SenseVoice-small</td><td align="center">234</td><td align="center">21.08</td><td align="center"><ins>6.52</ins></td><td align="center">8.05</td><td align="center"><b>7.34</b></td><td align="center">6.34</td><td align="center">5.74</td><td align="center"><ins>6.65</ins></td><td align="center">6.69</td><td align="center">9.95</td>
|
31 |
+
<tr>
|
32 |
+
<td align="left"><b>SenseVoice-s-Yue⭐</b></td><td align="center">234</td><td align="center">19.19</td><td align="center">6.71</td><td align="center">6.87</td><td align="center">8.68</td><td align="center"><ins>5.43</ins></td><td align="center">5.24</td><td align="center">6.93</td><td align="center">5.23</td><td align="center">8.63</td>
|
33 |
+
</tr>
|
34 |
+
</tr>
|
35 |
+
<tr>
|
36 |
+
<td align="left">Dolphin-small</td><td align="center">372</td><td align="center">59.20</td><td align="center">7.38</td><td align="center">39.69</td><td align="center">51.29</td><td align="center">26.39</td><td align="center">7.21</td><td align="center">9.68</td><td align="center">32.32</td><td align="center">58.20</td>
|
37 |
+
</tr>
|
38 |
+
<tr>
|
39 |
+
<td align="left">TeleASR</td><td align="center">700</td><td align="center">37.18</td><td align="center">7.27</td><td align="center">7.02</td><td align="center"><ins>7.88</ins></td><td align="center">6.25</td><td align="center">8.02</td><td align="center"><b>5.98</b></td><td align="center">6.23</td><td align="center">11.33</td>
|
40 |
+
</tr>
|
41 |
+
<tr>
|
42 |
+
<td align="left">Whisper-medium</td><td align="center">769</td><td align="center">75.50</td><td align="center">68.69</td><td align="center">59.44</td><td align="center">62.50</td><td align="center">62.31</td><td align="center">64.41</td><td align="center">80.41</td><td align="center">80.82</td><td align="center">50.96</td>
|
43 |
+
</tr>
|
44 |
+
<tr>
|
45 |
+
<td align="left"><b>Whisper-m-Yue⭐</b></td><td align="center">769</td><td align="center">18.69</td><td align="center">6.86</td><td align="center"><ins>6.86</ins></td><td align="center">11.03</td><td align="center">5.49</td><td align="center"><ins>4.70</ins></td><td align="center">8.51</td><td align="center"><ins>5.05</ins></td><td align="center"><ins>8.05</ins></td>
|
46 |
+
</tr>
|
47 |
+
|
48 |
+
<tr>
|
49 |
+
<td align="left">FireRedASR-AED-L</td><td align="center">1100</td><td align="center">73.70</td><td align="center">18.72</td><td align="center">43.93</td><td align="center">43.33</td><td align="center">34.53</td><td align="center">48.05</td><td align="center">49.99</td><td align="center">55.37</td><td align="center">50.26</td>
|
50 |
+
</tr>
|
51 |
+
<tr>
|
52 |
+
<td align="left">Whisper-large-v3</td><td align="center">1550</td><td align="center">45.09</td><td align="center">15.46</td><td align="center">12.85</td><td align="center">16.36</td><td align="center">14.63</td><td align="center">17.84</td><td align="center">20.70</td><td align="center">12.95</td><td align="center">26.86</td>
|
53 |
+
</tr>
|
54 |
+
|
55 |
+
<tr><td align="left" colspan="11"><b>w/ LLM</b></td></tr>
|
56 |
+
|
57 |
+
<tr>
|
58 |
+
<td align="left">Qwen2.5-Omni-3B</td><td align="center">3000</td><td align="center">72.01</td><td align="center">7.49</td><td align="center">12.59</td><td align="center">11.75</td><td align="center">38.91</td><td align="center">10.59</td><td align="center">25.78</td><td align="center">67.95</td><td align="center">88.46</td>
|
59 |
+
</tr>
|
60 |
+
<tr>
|
61 |
+
<td align="left">Kimi-Audio</td><td align="center">7000</td><td align="center">68.65</td><td align="center">24.34</td><td align="center">40.90</td><td align="center">38.72</td><td align="center">30.72</td><td align="center">44.29</td><td align="center">45.54</td><td align="center">50.86</td><td align="center">33.49</td>
|
62 |
+
</tr>
|
63 |
+
<tr>
|
64 |
+
<td align="left">FireRedASR-LLM-L</td><td align="center">8300</td><td align="center">73.70</td><td align="center">18.72</td><td align="center">43.93</td><td align="center">43.33</td><td align="center">34.53</td><td align="center">48.05</td><td align="center">49.99</td><td align="center">49.87</td><td align="center">45.92</td>
|
65 |
+
</tr>
|
66 |
+
<tr>
|
67 |
+
<td align="left"><b>Conformer-LLM-Yue⭐</b></td><td align="center">4200</td><td align="center"><ins>17.22</ins></td><td align="center"><b>6.21</b></td><td align="center"><b>6.23</b></td><td align="center">9.52</td><td align="center"><b>4.35</b></td><td align="center"><b>4.57</b></td><td align="center">6.98</td><td align="center"><b>4.73</b></td><td align="center"><b>7.91</b></td>
|
68 |
+
</tr>
|
69 |
+
</table>
|