@article{towardsholisticlanguagevideo, title = {Towards Holistic Language-video Representation: the language model-enhanced MSR-Video to Text Dataset}, author = {Yuchen Yang and Yingxuan Duan}, year = {2024}, eprint = {2406.13809}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2406.13809v1}, }