@article{videoopdefficientposttrainingofmultimoda, title = {Video-OPD: Efficient Post-Training of Multimodal Large Language Models for Temporal Video Grounding via On-Policy Distillation}, author = {Jiaze Li and Hao Yin and Haoran Xu and Boshen Xu and Wenhui Tan and Zewen He and Jianzhong Ju and Zhenbo Luo and Jian Luan}, year = {2026}, eprint = {2602.02994}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2602.02994}, }