@article{mementosacomprehensivebenchmarkfor, title = {Mementos: A Comprehensive Benchmark for Multimodal Large Language Model Reasoning over Image Sequences}, author = {Xiyao Wang and YuHang Zhou and Xiaoyu Liu and Hongjin Lu and Yuancheng Xu and Feihong He and Jaehong Yoon and Taixi Lu and Gedas Bertasius and Mohit Bansal and Huaxiu Yao and Furong Huang}, year = {2024}, eprint = {2401.10529}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2401.10529v2}, }