嘉渊
2023-05-12 8167ecaa49108c420c16deabdc4bdca3b75cf286
update repo
2个文件已修改
7 ■■■■ 已修改文件
funasr/datasets/large_datasets/dataset.py 3 ●●●● 补丁 | 查看 | 原始文档 | blame | 历史
funasr/fileio/sound_scp.py 4 ●●● 补丁 | 查看 | 原始文档 | blame | 历史
funasr/datasets/large_datasets/dataset.py
@@ -135,7 +135,8 @@
                            speed = random.choice(self.speed_perturb)
                            if speed != 1.0:
                                mat, _ = torchaudio.sox_effects.apply_effects_tensor(
                                    mat, sampling_rate, [['speed', str(speed)], ['rate', str(sampling_rate)]])
                                    torch.tensor(mat).view(1, -1), sampling_rate, [['speed', str(speed)], ['rate', str(sampling_rate)]])
                                mat = mat.view(-1).numpy()
                        sample_dict[data_name] = mat
                        sample_dict["sampling_rate"] = sampling_rate
                        if data_name == "speech":
funasr/fileio/sound_scp.py
@@ -8,6 +8,7 @@
import librosa
from typeguard import check_argument_types
import torch
import torchaudio
from funasr.fileio.read_text import read_2column_text
@@ -62,8 +63,9 @@
            speed = random.choice(self.speed_perturb)
            if speed != 1.0:
                array, _ = torchaudio.sox_effects.apply_effects_tensor(
                    array, rate,
                    torch.tensor(array).view(1, -1), rate,
                    [['speed', str(speed)], ['rate', str(rate)]])
            array = array.view(-1).numpy()
        return rate, array