--- dataset_info: features: - name: audio dtype: audio - name: transcription dtype: string - name: duration dtype: float64 - name: sr dtype: int64 - name: wav2vec2pred dtype: string - name: wer dtype: float64 - name: wps dtype: float64 - name: is_better dtype: bool splits: - name: commonvoice num_bytes: 26613419533.408 num_examples: 963636 download_size: 58948504311 dataset_size: 61107679180.05001 configs: - config_name: default data_files: - split: commonvoice path: data/commonvoice-* task_categories: - automatic-speech-recognition language: - bn size_categories: - 1M