python/FunASR-XL.git

			@@ -13,7 +13,7 @@
			sr = 8000
			n_mels = 23
			mel_basis = librosa.filters.mel(sr, n_fft, n_mels)
			Y = np.dot(Y ** 2, mel_basis.T)
			Y = np.dot(Y**2, mel_basis.T)
			Y = np.log10(np.maximum(Y, 1e-10))
			mean = np.mean(Y, axis=0)
			Y = Y - mean
			@@ -27,25 +27,21 @@


			def splice(Y, context_size=0):
			Y_pad = np.pad(
			Y,
			[(context_size, context_size), (0, 0)],
			'constant')
			Y_pad = np.pad(Y, [(context_size, context_size), (0, 0)], "constant")
			Y_spliced = np.lib.stride_tricks.as_strided(
			np.ascontiguousarray(Y_pad),
			(Y.shape[0], Y.shape[1] * (2 * context_size + 1)),
			(Y.itemsize * Y.shape[1], Y.itemsize), writeable=False)
			(Y.itemsize * Y.shape[1], Y.itemsize),
			writeable=False,
			)
			return Y_spliced


			def stft(
			data,
			frame_size=1024,
			frame_shift=256):
			def stft(data, frame_size=1024, frame_shift=256):
			fft_size = 1 << (frame_size - 1).bit_length()
			if len(data) % frame_shift == 0:
			return librosa.stft(data, n_fft=fft_size, win_length=frame_size,
			hop_length=frame_shift).T[:-1]
			return librosa.stft(data, n_fft=fft_size, win_length=frame_size, hop_length=frame_shift).T[
			:-1
			]
			else:
			return librosa.stft(data, n_fft=fft_size, win_length=frame_size,
			hop_length=frame_shift).T
			return librosa.stft(data, n_fft=fft_size, win_length=frame_size, hop_length=frame_shift).T