zhifu gao
2023-04-17 2e81f55ecca9172d8f3349518437049e955894d3
funasr/datasets/large_datasets/utils/tokenize.py
@@ -47,8 +47,8 @@
    length = len(text)
    for i in range(length):
        x = text[i]
        if i == length-1 and "punc" in data and text[i].startswith("vad:"):
            vad = x[-1][4:]
        if i == length-1 and "punc" in data and x.startswith("vad:"):
            vad = x[4:]
            if len(vad) == 0:
                vad = -1
            else: