nichongjia-2007
2023-05-31 cc2c1d1d53dea5d2c45f858d1baa5bd279f47987
egs/aishell/paraformer/conf/train_asr_paraformer_conformer_12e_6d_2048_256.yaml
@@ -29,6 +29,17 @@
    self_attention_dropout_rate: 0.0
    src_attention_dropout_rate: 0.0
# frontend related
frontend: wav_frontend
frontend_conf:
    fs: 16000
    window: hamming
    n_mels: 80
    frame_length: 25
    frame_shift: 10
    lfr_m: 1
    lfr_n: 1
model: paraformer
model_conf:
    ctc_weight: 0.3
@@ -36,16 +47,12 @@
    length_normalized_loss: false
    predictor_weight: 1.0
    sampling_ratio: 0.4
# minibatch related
batch_type: length
batch_bins: 25000
num_workers: 16
    use_1st_decoder_loss: true
# optimization related
accum_grad: 1
grad_clip: 5
max_epoch: 50
max_epoch: 150
val_scheduler_criterion:
    - valid
    - acc
@@ -78,7 +85,7 @@
    - 40
    num_time_mask: 2
predictor: cif_predictor_v2
predictor: cif_predictor
predictor_conf:
    idim: 256
    threshold: 1.0
@@ -86,6 +93,17 @@
    r_order: 1
    tail_threshold: 0.45
dataset_conf:
    data_names: speech,text
    data_types: sound,text
    shuffle: True
    shuffle_conf:
        shuffle_size: 2048
        sort_size: 500
    batch_conf:
        batch_type: token
        batch_size: 25000
    num_workers: 8
log_interval: 50
normalize: None