RAD-MMM
A TTS model that makes a speaker speak new languages
File Explorer
Download Latest Version (.zip)- opensource_data_16khz.yaml
- RADMMM_16khz_model_config.yaml
- RADMMM_durationmodel_config.yaml
- RADMMM_energymodel_config.yaml
- RADMMM_epic_durationmodel_config.yaml
- RADMMM_f0model_config.yaml
- RADMMM_LJS_22khz_data_config.yaml
- RADMMM_LJS_data_config.yaml
- RADMMM_model_config.yaml
- RADMMM_opensource_16khz_data_config.yaml
- RADMMM_opensource_data_config_phonemizerless.yaml
- RADMMM_train_config.yaml
- RADMMM_vpredmodel_config.yaml
- RADTTS_durationmodel_config.yaml
- RADTTS_energymodel_config.yaml
- RADTTS_epic_durationmodel_config.yaml
- RADTTS_f0model_config.yaml
- RADTTS_model_config.yaml
- RADTTS_vpredmodel_config.yaml
- berndungerer_audiopath_text_sid_emotion_duration_train_filelist_filtered.txt
- berndungerer_audiopath_text_sid_emotion_duration_train_filelist_filtered_phonemized.txt
- berndungerer_audiopath_text_sid_emotion_duration_val_filelist.txt
- berndungerer_audiopath_text_sid_emotion_duration_val_filelist_phonemized.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration_test.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration_train.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration_train_phonemized.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration_val.txt
- hi_indic_iiit_hyderbad_audiopath_text_sid_emotion_duration_val_phonemized.txt
- ljs_audiopath_text_sid_emotion_duration_test_filelist.txt
- ljs_audiopath_text_sid_emotion_duration_train_filelist.txt
- ljs_audiopath_text_sid_emotion_duration_train_filelist_phonemized.txt
- ljs_audiopath_text_sid_emotion_duration_val_filelist.txt
- ljs_audiopath_text_sid_emotion_duration_val_filelist_phonemized.txt
- tux_audiopath_text_sid_emotion_duration_train_filelist_filtered.txt
- tux_audiopath_text_sid_emotion_duration_train_filelist_filtered_phonemized.txt
- tux_audiopath_text_sid_emotion_duration_val_filelist.txt
- tux_audiopath_text_sid_emotion_duration_val_filelist_phonemized.txt
- ks_audiopath_text_sid_emotion_duration_filelist.txt
- ks_audiopath_text_sid_emotion_duration_train_filelist.txt
- ks_audiopath_text_sid_emotion_duration_train_filelist_phonemized.txt
- ks_audiopath_text_sid_emotion_duration_val_filelist.txt
- ks_audiopath_text_sid_emotion_duration_val_filelist_phonemized.txt
- nadineeckert_audiopath_text_sid_emotion_duration_train_filelist_filtered.txt
- nadineeckert_audiopath_text_sid_emotion_duration_train_filelist_filtered_phonemized.txt
- nadineeckert_audiopath_text_sid_emotion_duration_val_filelist.txt
- nadineeckert_audiopath_text_sid_emotion_duration_val_filelist_phonemized.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration_test.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration_train.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration_train_phonemized.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration_val.txt
- ed_portuguese_audiopath_transcript_sid_emotion_duration_val_phonemized.txt
- ed_portuguese_audiopath_transcript_sid_test.txt
- ed_portuguese_audiopath_transcript_sid_train.txt
- ed_portuguese_audiopath_transcript_sid_val.txt
- collated_stats.json
- Hindi_F-other.json
- Hindi_M-other.json
- Marathi_F-other.json
- Marathi_M-other.json
- opensource_collated_stats.json
- Telugu_F-other.json
- Telugu_M-other.json
- 22khz-limmits-nonparallel-processed.json
- 22khz-limmits-nonparallel.json
- 22khz-limmits-parallel-processed.json
- 22khz-limmits-parallel.json
- 22khz-ljs.json
- real_22khz_ljs.json
- Dockerfile
- language_transfer_prompts.json
- resynthesis_prompts.json
- radmmm.py
- compute_speaker_prosody_statistics.py
- scripting_utils.py
- abbreviations.py
- acronyms.py
- cleaners.py
- cmudict-0.7b
- cmudict.py
- datestime.py
- grapheme_dictionary.py
- heteronyms
- letters_and_numbers.py
- LICENSE
- numerical.py
- symbols.py
- text_processing.py
- radmmm_data_table.png
- ljs_audio_text_test_filelist.txt
- ljs_audio_text_train_filelist.txt
- ljs_audio_text_val_filelist.txt
- __init__.py
- cleaners.py
- cmudict.py
- LICENSE
- numbers.py
- symbols.py
- ljs_audio_text_test_filelist.txt
- ljs_audio_text_train_filelist.txt
- ljs_audio_text_val_filelist.txt
- __init__.py
- cleaners.py
- cmudict.py
- LICENSE
- numbers.py
- symbols.py
- audio_processing.py
- data_utils.py
- demo.wav
- distributed.py
- Dockerfile
- fp16_optimizer.py
- hparams.py
- inference.ipynb
- layers.py
- LICENSE
- logger.py
- loss_function.py
- loss_scaler.py
- model.py
- multiproc.py
- plotting_utils.py
- README.md
- requirements.txt
- stft.py
- tensorboard.png
- train.py
- utils.py
- .gitmodules
- config.json
- convert_model.py
- denoiser.py
- distributed.py
- glow.py
- glow_old.py
- inference.py
- LICENSE
- mel2samp.py
- README.md
- requirements.txt
- train.py
- waveglow_logo.png
- .gitmodules
- audio_processing.py
- data_utils.py
- demo.wav
- distributed.py
- Dockerfile
- hparams.py
- inference.ipynb
- layers.py
- LICENSE
- logger.py
- loss_function.py
- loss_scaler.py
- model.py
- multiproc.py
- plotting_utils.py
- README.md
- requirements.txt
- stft.py
- tensorboard.png
- train.py
- utils.py
- .gitmodules
- config.json
- convert_model.py
- denoiser.py
- distributed.py
- glow.py
- glow_old.py
- infer_indic.py
- infer_ljs.py
- inference.py
- LICENSE
- mel2samp.py
- README.md
- requirements.txt
- train.py
- waveglow_logo.png
- hifigan_denoiser.py
- hifigan_env.py
- hifigan_models.py
- hifigan_utils.py
- vocoder_utils.py
- .gitignore
- alignment.py
- attribute_predictors.py
- audio_processing.py
- common.py
- data.py
- data_modules.py
- decoders.py
- inference.ipynb
- LICENSE
- loss.py
- maskedbatchnorm1d.py
- partialconv1d.py
- plotting_utils.py
- radam.py
- README.md
- requirements.txt
- splines.py
- stft_loss.py
- training_callbacks.py
- tts_lightning_modules.py
- tts_main.py
- utils.py
- wave_transforms.py
// repository documentation
Was this content helpful?
(0 ratings)
