From 3c349ac0531b07239f37b81254f8568ab80e3f6a Mon Sep 17 00:00:00 2001
From: Han Zhang <45134013+holazzer@users.noreply.github.com>
Date: 星期二, 18 三月 2025 11:45:37 +0800
Subject: [PATCH] fix: use converted token_ids for alignment for sensevoice model with timestamp output (#2429)
---
examples/industrial_data_pretraining/fsmn_vad_streaming/demo.py | 4 +---
1 files changed, 1 insertions(+), 3 deletions(-)
diff --git a/examples/industrial_data_pretraining/fsmn_vad_streaming/demo.py b/examples/industrial_data_pretraining/fsmn_vad_streaming/demo.py
index 0f30a37..21ce0cb 100644
--- a/examples/industrial_data_pretraining/fsmn_vad_streaming/demo.py
+++ b/examples/industrial_data_pretraining/fsmn_vad_streaming/demo.py
@@ -9,11 +9,9 @@
model = AutoModel(model="iic/speech_fsmn_vad_zh-cn-16k-common-pytorch")
-mm = model.model
-for p in mm.parameters():
- print(f"{p.numel()}")
res = model.generate(input=wav_file)
print(res)
+
# [[beg1, end1], [beg2, end2], .., [begN, endN]]
# beg/end: ms
--
Gitblit v1.9.1