@article{unimo3multigranularityinteractionfor, title = {UNIMO-3: Multi-granularity Interaction for Vision-Language Representation Learning}, author = {Hao Yang and Can Gao and Hao Líu and Xinyan Xiao and Yanyan Zhao and Bing Qin}, year = {2023}, eprint = {2305.13697}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2305.13697v1}, }