zhifu gao
2024-05-08 b1c186fd00fef54bcad3aa1d073a1a313642d641
funasr/datasets/sense_voice_datasets/datasets.py
@@ -112,7 +112,7 @@
            eos = self.tokenizer.encode(self.eos, allowed_special="all")  # [eos]
            ids = prompt_ids + target_ids + eos
            ids = prompt_ids + target_ids + eos  # [sos, task, lid, text, eos]
            ids_lengths = len(ids)
            text = torch.tensor(ids, dtype=torch.int64)