@article{temporalconditionalreferringvideoobjects, title = {Temporal-Conditional Referring Video Object Segmentation with Noise-Free Text-to-Video Diffusion Model}, author = {Ruixin Zhang and Jiaqing Fan and Yifan Liao and Qian Qiao and Fanzhang Li}, year = {2025}, eprint = {2508.13584}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2508.13584}, }