VeritasPP
No description available.
File Explorer
- banner.png
- ding.png
- discord_qr.jpg
- wechat.png
- grpo.png
- megatron.png
- npu.png
- asyncengine.png
- deepeyes.png
- dpo_data.png
- grpo.png
- grpo_clevr_count.png
- grpo_code.png
- grpo_countdown.png
- grpo_countdown_1.png
- grpo_geoqa.png
- grpo_multi_turn.png
- grpo_openr1_multimodal.png
- gym_env.png
- kto_data.png
- multiturn_pipeline.png
- sapo_tau.png
- treepo.png
- web-ui-en.jpg
- web-ui.jpg
- class.rst
- classtemplate.rst
- sobolengine.rst
- Elastic.md
- Embedding.md
- GRPO-Code-Training.md
- GRPO-Multi-Modal-Training.md
- GRPO.md
- Metax-support.md
- MLLM-Registration.md
- More-Best-Practices.md
- NPU-support.md
- Qwen3-Best-Practice.md
- Qwen3-VL-Best-Practice.md
- Qwen3_5-Best-Practice.md
- Rapidly-Training-VL-model.md
- Reranker.md
- Architecture.md
- Custom-dataset.md
- Custom-model.md
- Quick-start.md
- SWIFT-installation.md
- Web-UI.md
- CHORD.md
- CISPO.md
- DAPO.md
- deepeyes.md
- entropy_mask.md
- GSPO.md
- index.rst
- REINFORCEPP.md
- RLOO.md
- SAPO.md
- training_inference_mismatch.md
- treepo.md
- gym_env.md
- index.rst
- loss_types.md
- multi_task.md
- multi_turn.md
- reward_function.md
- reward_model.md
- GRPO.md
- index.rst
- index.rst
- Agent-support.md
- Command-line-parameters.md
- Evaluation.md
- Export-and-push.md
- Frequently-asked-questions.md
- GKD.md
- Inference-and-deployment.md
- Pre-training-and-Fine-tuning.md
- Ray.md
- Reinforced-Fine-tuning.md
- RLHF.md
- Sample.md
- Supported-models-and-datasets.md
- Use-tuners.md
- Ascend.md
- Command-line-parameters.md
- GKD.md
- GRPO.md
- LoRA-Training.md
- Mcore-Bridge.md
- Multimodal-Model.md
- Quick-start.md
- .readthedocs.yaml
- conf.py
- index.rst
- class.rst
- classtemplate.rst
- sobolengine.rst
- Elastic.md
- Embedding.md
- GRPO-Code-Training.md
- GRPO-Multi-Modal-Training.md
- GRPO.md
- Metax-support.md
- MLLM-Registration.md
- More-Best-Practices.md
- NPU-support.md
- Qwen3-Best-Practice.md
- Qwen3-VL-Best-Practice.md
- Qwen3_5-Best-Practice.md
- Rapidly-Training-VL-model.md
- Reranker.md
- Architecture.md
- Custom-dataset.md
- Custom-model.md
- Quick-start.md
- SWIFT-installation.md
- Web-UI.md
- CHORD.md
- CISPO.md
- DAPO.md
- deepeyes.md
- entropy_mask.md
- GSPO.md
- index.rst
- REINFORCEPP.md
- RLOO.md
- SAPO.md
- training_inference_mismatch.md
- treepo.md
- gym_env.md
- index.rst
- loss_types.md
- multi_task.md
- multi_turn.md
- reward_function.md
- reward_model.md
- GRPO.md
- index.rst
- index.rst
- Agent-support.md
- Command-line-parameters.md
- Evaluation.md
- Export-and-push.md
- Frequently-asked-questions.md
- GKD.md
- Inference-and-deployment.md
- Pre-training-and-Fine-tuning.md
- Ray.md
- Reinforced-Fine-tuning.md
- RLHF.md
- Sample.md
- Supported-models-and-datasets.md
- Use-tuners.md
- Ascend.md
- Command-line-parameters.md
- GKD.md
- GRPO.md
- LoRA-Training.md
- Mcore-Bridge.md
- Multimodal-Model.md
- Quick-start.md
- .readthedocs.yaml
- conf.py
- index.rst
- make.bat
- Makefile
- README.md
- demo.py
- demo.sh
- sglang.sh
- vllm.sh
- mllm.sh
- fsdp2.json
- train.sh
- vllm.sh
- dp_tp.sh
- train_sft_full.sh
- node1.sh
- node2.sh
- fsdp.json
- train.sh
- qwen3_lora_deepspeed.sh
- qwen3_lora_megatron.sh
- qwen3_next_megatron.sh
- qwen3_omni_full_mindspeed.sh
- moe_full_mindspeed.sh
- my_register.py
- test_register.py
- train.py
- dataset.py
- infer.sh
- model.py
- model_hf.py
- sft.sh
- client.py
- server.sh
- client.py
- server.sh
- openai_client.py
- swift_client.py
- openai_client.py
- swift_client.py
- openai_client.py
- swift_client.py
- client.py
- server.sh
- client.py
- server.sh
- client.py
- client_generative.py
- server.sh
- client.py
- server.sh
- client.py
- server.sh
- README.md
- sglang.sh
- vllm.sh
- vllm_dp.sh
- demo.py
- eval.sh
- sglang.sh
- vllm.sh
- train.sh
- eval.sh
- bnb.sh
- gptq.sh
- awq.sh
- bnb.sh
- fp8.sh
- gptq.sh
- awq.sh
- bnb.sh
- fp8.sh
- gptq.sh
- gptq.sh
- bnb.sh
- gptq.sh
- awq.sh
- bnb.sh
- fp8.sh
- gptq.sh
- gptq_v2.sh
- merge_lora.sh
- ollama.sh
- push_to_hub.sh
- batch_ddp.sh
- mllm_tp.sh
- distill_qwen3_235b.sh
- mtp.sh
- tp.sh
- batch_ddp.sh
- bert.sh
- lora.sh
- mllm_device_map.sh
- prm.sh
- reward_model.sh
- dp_tp.sh
- mllm_ddp.sh
- mllm_tp.sh
- mtp.sh
- cli_demo.sh
- demo.py
- demo_mllm.py
- demo_reward_model.py
- demo_vllm_reasoning_parser.py
- deepspeed.sh
- 72b_offload.sh
- qwen3_32b.sh
- qwen3_emb.sh
- qwen3_vl_emb.sh
- full.sh
- lora.sh
- benchmark.sh
- llm.sh
- vlm.sh
- dense_colocate.sh
- dense_server.sh
- moe_colocate_full.sh
- moe_colocate_lora.sh
- sapo.sh
- dense.sh
- dpo.sh
- loss_scale.sh
- moe.sh
- mtp.sh
- new_special_tokens.sh
- qwen3_235b.sh
- dense.sh
- moe.sh
- moe.sh
- new_special_tokens.sh
- seq_cls.sh
- deepseek_v3.sh
- moe.sh
- qwen3_moe.sh
- qwen3_moe_offload.sh
- node1.sh
- node2.sh
- dpo.sh
- full.sh
- lora.sh
- sft.sh
- full_dpo_offload.sh
- lora.sh
- dense.sh
- moe.sh
- qwen3_reranker.sh
- qwen3_vl_reranker.sh
- dense.sh
- group_by_length.sh
- moe.sh
- packing.sh
- dense.sh
- opsd.sh
- teacher_server.sh
- dense.sh
- moe.sh
- dense.sh
- moe.sh
- infer.sh
- train.sh
- full.sh
- base_to_chat.sh
- long_text.sh
- muon.sh
- pretrain.sh
- sft.sh
- infer.py
- train.sh
- train.sh
- train.sh
- flash.sh
- mcore.sh
- internvl3_5_gpt.sh
- mcore.sh
- train.sh
- train.sh
- train.sh
- train.sh
- mcore.sh
- train.sh
- train.sh
- mcore.sh
- mcore_grpo_moe.sh
- packing.sh
- transformers.sh
- mcore.sh
- mtp.sh
- non_padding_free.sh
- transformers.sh
- transformers.sh
- zero3.sh
- mcore.sh
- mcore_full.sh
- mixed.sh
- transformers.sh
- zero3.sh
- infer.ipynb
- infer.sh
- self-cognition-sft.ipynb
- sft.sh
- zh.ipynb
- infer.ipynb
- ocr-sft.ipynb
- distill.sh
- distill.yaml
- sample.sh
- sampling.yaml
- infer_lora.py
- train.sh
- deepseek_r1.sh
- glm4.sh
- qwen2_5.sh
- infer.sh
- train.sh
- full.sh
- lora.sh
- lora2.sh
- dpo.sh
- mcore.sh
- pretrained.sh
- reranker.sh
- seq_cls.sh
- sft.sh
- vlm.sh
- lora_sft.sh
- infer.py
- qwen3_emb.sh
- qwen3_vl_emb.sh
- train_gme.sh
- mcore.sh
- transformers.sh
- dft.sh
- infer.sh
- qwen2_5_32b.sh
- train.sh
- agent.sh
- grpo_32b_full.sh
- grpo_7b.sh
- moe_full.sh
- moe_lora.sh
- README.md
- vllm_gym.sh
- vllm_multi_turn.sh
- chord.sh
- full_lmdeploy.sh
- gspo.sh
- moe_full.sh
- moe_lora.sh
- qlora.sh
- README.md
- reinforce_plus_plus.sh
- rloo.sh
- sapo.sh
- transformers.sh
- vllm_72b_4gpu.sh
- vllm_lora_qwenvl72b.sh
- vllm_multi_turn.sh
- vllm_vl7b.sh
- colocate_multi_node1.sh
- colocate_multi_node2.sh
- Qwen2_5_32B_full.sh
- server_multi_node.sh
- train_dlc.sh
- deepeyes.sh
- deepeyes_plugin.py
- gsm8k.sh
- gsm8k_plugin.py
- tree_rollout.py
- tree_rollout.sh
- tree_rollout_plugin.py
- plugin.py
- run_external_reward_func.sh
- run_external_reward_model.sh
- run_external_scheduler.sh
- grpo.sh
- infer.sh
- prompt.txt
- sft.sh
- llama4.sh
- qwen3_moe.sh
- train.sh
- train.sh
- train_zero2.sh
- train_zero3.sh
- train.sh
- fsdp2.json
- train.sh
- fsdp_offload.json
- train.sh
- multi_node.yaml
- train_node1.sh
- train_node2.sh
- host.txt
- README.md
- train.sh
- train.sh
- sft.sh
- sft.yaml
- train_node1.sh
- train_node2.sh
- train_node1.sh
- train_node2.sh
- infer.sh
- merge_lora.sh
- seq_cls.sh
- sft.sh
- infer.sh
- sft.sh
- full.sh
- lora.sh
- fast.sh
- full.sh
- kto.sh
- audio.sh
- caption.sh
- grounding.sh
- infer.sh
- ocr.sh
- video.sh
- vit_gradient_checkpointing.sh
- infer.sh
- merge_lora.sh
- tokens.txt
- train.sh
- muon.sh
- muonclip.sh
- dpo.sh
- dpo_vlm.sh
- liger_kernel.sh
- llm.sh
- streaming.sh
- dpo_vlm.sh
- sft.sh
- loss_scale.sh
- tuner_phi4_mm.sh
- train.sh
- train.sh
- merge_lora.sh
- train.sh
- merge_lora.sh
- train.sh
- gptq.sh
- hqq.sh
- infer.py
- train_generative_reranker.sh
- train_generative_reranker_listwise.sh
- train_reranker.sh
- train_reranker_auto_patch.sh
- train_reranker_listwise.sh
- train_reranker_mm.sh
- math.json
- rft.py
- full.sh
- lora.sh
- fast.sh
- full.sh
- teacher_server.sh
- think_model.sh
- vllm_colocate.sh
- vllm_server.sh
- opsd.sh
- full.sh
- lora.sh
- cpo.sh
- kto.sh
- mpo.sh
- README.md
- rm.sh
- simpo.sh
- deploy.sh
- infer.sh
- sft.sh
- infer.py
- infer.sh
- sft.sh
- vlm.sh
- deploy.sh
- infer.sh
- sft.sh
- infer.py
- infer.sh
- sft.sh
- deploy.sh
- infer.sh
- sft.sh
- sequence_parallel.sh
- sequence_parallel_512k.sh
- sequence_parallel_dpo.sh
- sequence_parallel_emb.sh
- sequence_parallel_grpo.sh
- sequence_parallel_reranker.sh
- sequence_parallel_seq_cls.sh
- lazy_tokenize.sh
- streaming.sh
- deepseek_r1.sh
- qwen3_demo1.sh
- qwen3_demo2.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train_galore.sh
- train_qgalore.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- train.sh
- infer.sh
- lora_sft.sh
- on_policy_distillation.sh
- infer.sh
- infer.yaml
- sft.json
- sft.sh
- sft.yaml
- infer.sh
- infer.yaml
- sft.sh
- sft.yaml
- README.md
- method_overview2.png
- vaopd.png
- dependency_links.txt
- entry_points.txt
- not-zip-safe
- PKG-INFO
- requires.txt
- SOURCES.txt
- top_level.txt
- docs.txt
- eval.txt
- framework.txt
- install_all.sh
- ray.txt
- swanlab.txt
- tests.txt
- tuner.json
- exp.py
- exp_utils.py
- generate_report.py
- plot_loss.py
- run_dataset_info.py
- run_model_info.py
- run_template.py
- test_link_valid.py
- deploy_model.sh
- infer_vllm_single.py
- __init__.py
- base.py
- deepseek_v3_1.py
- extra.py
- glm4.py
- hermes.py
- llama.py
- mapping.py
- minimax_m2.py
- mistral.py
- qwen.py
- qwen3_coder.py
- react.py
- seed_oss.py
- toolbench.py
- youtu.py
- __init__.py
- base_args.py
- data_args.py
- generation_args.py
- model_args.py
- quant_args.py
- template_args.py
- __init__.py
- app_args.py
- deploy_args.py
- eval_args.py
- export_args.py
- infer_args.py
- merge_args.py
- pretrain_args.py
- sft_args.py
- tuner_args.py
- webui_args.py
- __init__.py
- activation_cpu_offload.py
- adalora.py
- base.py
- deepspeed_elastic.py
- early_stop.py
- lisa.py
- mapping.py
- perf_log.py
- __init__.py
- export.py
- main.py
- pt.py
- sft.py
- __init__.py
- app.py
- deploy.py
- eval.py
- export.py
- infer.py
- main.py
- merge_lora.py
- pt.py
- sft.py
- utils.py
- web_ui.py
- fsdp2.json
- zero0.json
- zero1.json
- zero2.json
- zero2_offload.json
- zero3.json
- zero3_offload.json
- __init__.py
- dispatcher.py
- shard.py
- dataset_info.json
- __init__.py
- llm.py
- mllm.py
- __init__.py
- core.py
- extra.py
- __init__.py
- dataset_meta.py
- dataset_syntax.py
- indexed_dataset.py
- loader.py
- media.py
- packing.py
- register.py
- utils.py
- __init__.py
- constant.py
- hub.py
- __init__.py
- base.py
- grpo_vllm_engine.py
- infer_client.py
- infer_engine.py
- lmdeploy_engine.py
- patch.py
- protocol.py
- sglang_engine.py
- transformers_engine.py
- utils.py
- vllm_engine.py
- __init__.py
- base.py
- causal_lm.py
- embedding.py
- mapping.py
- reranker.py
- agentflan.json
- alpha_umi.json
- hermes.json
- ignore_empty_think.json
- qwen.json
- react.json
- __init__.py
- agent.py
- base.py
- mapping.py
- other.py
- utils.py
- __init__.py
- export_args.py
- megatron_args.py
- megatron_base_args.py
- pretrain_args.py
- sft_args.py
- __init__.py
- base.py
- default_flow.py
- mapping.py
- print.py
- swanlab.py
- tensorboard.py
- utils.py
- wandb.py
- __init__.py
- utils.py
- __init__.py
- export.py
- __init__.py
- pretrain.py
- sft.py
- __init__.py
- __init__.py
- base.py
- batch_sampler.py
- dpo_trainer.py
- embedding_trainer.py
- gkd_trainer.py
- grpo_trainer.py
- kto_trainer.py
- reranker_trainer.py
- reward_trainer.py
- rlhf_mixin.py
- rollout_mixin.py
- trainer.py
- utils.py
- vocab_parallel_utils.py
- __init__.py
- convert_utils.py
- megatron_lm_utils.py
- parallel_utils.py
- patcher.py
- router_replay_utils.py
- utils.py
- __init__.py
- convert.py
- init.py
- __init__.py
- acc.py
- base.py
- embedding.py
- mapping.py
- nlg.py
- reranker.py
- utils.py
- __init__.py
- baai.py
- baichuan.py
- baidu.py
- bert.py
- codefuse.py
- deepseek.py
- gemma.py
- glm.py
- internlm.py
- llama.py
- llava.py
- llm.py
- mamba.py
- microsoft.py
- minicpm.py
- minimax.py
- mistral.py
- mllm.py
- moonshot.py
- mplug.py
- openbuddy.py
- qwen.py
- seed.py
- skywork.py
- stepfun.py
- telechat.py
- tencent.py
- valley.py
- yi.py
- __init__.py
- constant.py
- model_arch.py
- model_meta.py
- npu_patcher.py
- patcher.py
- register.py
- utils.py
- __init__.py
- adafactor.py
- adamw.py
- adamw8bit.py
- galore_projector.py
- utils.py
- __init__.py
- base.py
- lorap.py
- mapping.py
- multimodal.py
- muon.py
- muonclip.py
- __init__.py
- app.py
- build_ui.py
- locale.py
- __init__.py
- eval.py
- utils.py
- __init__.py
- cached_dataset.py
- export.py
- merge_lora.py
- ollama.py
- quant.py
- __init__.py
- deploy.py
- infer.py
- rollout.py
- utils.py
- __init__.py
- base.py
- distill_sampler.py
- sampling.py
- utils.py
- vanilla_sampler.py
- __init__.py
- kto.py
- pretrain.py
- rlhf.py
- sft.py
- tuner.py
- __init__.py
- base.py
- utils.py
- __init__.py
- arguments.py
- base.py
- resource_manager.py
- __init__.py
- orm.py
- prm.py
- rm_plugin.py
- __init__.py
- args_mixin.py
- arguments.py
- cpo_trainer.py
- dpo_trainer.py
- gkd_trainer.py
- grpo_trainer.py
- kto_trainer.py
- orpo_trainer.py
- ppo_trainer.py
- reward_trainer.py
- rlhf_mixin.py
- rollout_mixin.py
- utils.py
- vllm_client.py
- __init__.py
- gym_env.py
- multi_turn.py
- __init__.py
- ulysses.py
- utils.py
- zigzag_ring_attn.py
- __init__.py
- baai.py
- baidu.py
- bert.py
- deepseek.py
- dots.py
- gemma.py
- glm.py
- idefics3.py
- internlm.py
- internvl.py
- kwai.py
- llama.py
- llava.py
- llm.py
- megrez.py
- microsoft.py
- midashenglm.py
- minicpm.py
- minimax.py
- minimind.py
- mistral.py
- molmo.py
- moonshot.py
- mplug.py
- openbuddy.py
- pixtral.py
- qwen.py
- seed.py
- self_aug.py
- stepfun.py
- tencent.py
- utils.py
- valley.py
- yi.py
- __init__.py
- base.py
- constant.py
- grounding.py
- register.py
- template_inputs.py
- template_meta.py
- utils.py
- vision_utils.py
- __init__.py
- arguments.py
- embedding_trainer.py
- mixin.py
- patcher.py
- reranker_trainer.py
- seq2seq_trainer.py
- trainer.py
- trainer_factory.py
- utils.py
- __init__.py
- base.py
- dummy.py
- ia3.py
- lora_llm.py
- mapping.py
- __init__.py
- llama.py
- longlora.py
- __init__.py
- scetuning.py
- scetuning_components.py
- __init__.py
- adapter.py
- base.py
- llamapro.py
- lora.py
- lora_layers.py
- mapping.py
- neftune.py
- part.py
- peft.py
- prompt.py
- reft.py
- restuning.py
- restuning_components.py
- side.py
- utils.py
- __init__.py
- eval.py
- llm_eval.py
- model.py
- runtime.py
- __init__.py
- export.py
- llm_export.py
- model.py
- runtime.py
- __init__.py
- advanced.py
- dataset.py
- external_rollout.py
- external_runtime.py
- grpo_advanced.py
- hyper.py
- llm_grpo.py
- lora.py
- model.py
- optimizer.py
- quantization.py
- report_to.py
- reward.py
- rollout.py
- runtime.py
- save.py
- target.py
- tuner.py
- __init__.py
- generate.py
- llm_infer.py
- model.py
- runtime.py
- __init__.py
- advanced.py
- dataset.py
- hyper.py
- llm_rlhf.py
- lora.py
- model.py
- optimizer.py
- quantization.py
- report_to.py
- rlhf.py
- runtime.py
- save.py
- target.py
- tuner.py
- __init__.py
- llm_sample.py
- model.py
- runtime.py
- sample.py
- __init__.py
- advanced.py
- dataset.py
- hyper.py
- llm_train.py
- lora.py
- model.py
- optimizer.py
- quantization.py
- report_to.py
- runtime.py
- save.py
- self_cog.py
- target.py
- task.py
- tuner.py
- utils.py
- __init__.py
- app.py
- base.py
- __init__.py
- constants.py
- env.py
- hf_config.py
- hub_utils.py
- import_utils.py
- io_utils.py
- logger.py
- np_utils.py
- processor_utils.py
- shutdown_manager.py
- tb_utils.py
- transformers_utils.py
- utils.py
- __init__.py
- version.py
- test_app.py
- test_dataset.py
- test_logprobs.py
- test_eval.py
- test_quant.py
- test_arch.py
- test_dataset.py
- test_model.py
- test_stream.py
- test_template.py
- __init__.py
- test_check_model.py
- test_agent.py
- test_infer.py
- test_logprobs.py
- test_main.py
- test_max_memory.py
- test_mllm.py
- test_sglang.py
- infer.json
- sft.json
- alpaca.csv
- alpaca.jsonl
- alpaca2.csv
- chatml.jsonl
- conversations.jsonl
- multi_modal_1.jsonl
- multi_modal_2.jsonl
- multi_modal_3.jsonl
- sharegpt.jsonl
- swift_multi.json
- swift_multi.jsonl
- swift_pre.csv
- swift_pre.jsonl
- swift_single.csv
- swift_single.jsonl
- __init__.py
- test_custom.py
- test_dataset.py
- test_ollama_export.py
- test_run.py
- test_template.py
- test_export.py
- test_embedding.py
- test_export.py
- test_gkd.py
- test_grpo.py
- test_kto.py
- test_lora.py
- test_rlhf.py
- test_flash_attn.py
- test_llm.py
- test_mllm.py
- test_client.py
- test_agent.py
- test_audio.py
- test_gene.py
- test_llm.py
- test_template.py
- test_tool.py
- test_cls.py
- test_lmdeploy_vlm.py
- test_padding_side.py
- test_rlhf_loss.py
- test_channel.py
- test_cls.py
- test_embedding.py
- test_export_cached_dataset.py
- test_freeze.py
- test_gkd.py
- test_grounding.py
- test_grpo.py
- test_kto.py
- test_liger.py
- test_multilabel.py
- test_packing.py
- test_ppo.py
- test_pt.py
- test_resume_from_checkpoint.py
- test_rlhf.py
- test_sample.py
- test_sft.py
- test_train_eval.py
- test_vit_lr.py
- test_vllm_importance_sampling_basic.py
- __init__.py
- test_extra_state_dict.py
- test_merged_linear.py
- test_neft.py
- test_peft.py
- test_scetuning.py
- test_swift_base.py
- test_swift_device_map.py
- test_swift_restuning.py
- __init__.py
- test_async_rewards.py
- test_file_utils.py
- test_io_utils.py
- test_rewards.py
- test_split_str_parts_by.py
- test_torch_utils.py
- __init__.py
- model_tag.py
- run.py
- run_config.yaml
- LICENSE
- Makefile
- MANIFEST.in
- README.md
- requirements.txt
- setup.cfg
- setup.py
# Use via CDN
jsDelivrjsDelivr serves any public GitHub repository as a CDN with zero setup. Pick a version and a file to get a ready-to-paste link and snippet.
Link
Example
// repository documentation
Was this content helpful?
(0 ratings)
