From 4d907718f39e2b0f7a0c714c2e3de289e742fc61 Mon Sep 17 00:00:00 2001
From: Carl <415692979@qq.com>
Date: 星期四, 28 三月 2024 13:42:00 +0800
Subject: [PATCH] 修正commit 87b62d68957a2194b017a43b6c2a15424a05a984 引入的英文整句标点预测导致末尾两个单词中间的空格被删除的问题。 (#1556)

---
 funasr/datasets/audio_datasets/index_ds.py |    3 ++-
 1 files changed, 2 insertions(+), 1 deletions(-)

diff --git a/funasr/datasets/audio_datasets/index_ds.py b/funasr/datasets/audio_datasets/index_ds.py
index 12ffd23..34f7b4f 100644
--- a/funasr/datasets/audio_datasets/index_ds.py
+++ b/funasr/datasets/audio_datasets/index_ds.py
@@ -99,7 +99,8 @@
                     target = data["target"]
                     source_len = data.get("source_len", 1)
                     target_len = data.get("target_len", 0)
-                    
+                    if "aishell" in source:
+                        target = target.replace(" ", "")
                     contents.append({"source": source,
                                      "prompt": prompt,
                                      "target": target,

--
Gitblit v1.9.1