@article{fusiontoenhancefusionvisualencodertoenha, title = {Fusion to Enhance: Fusion Visual Encoder to Enhance Multimodal Language Model}, author = {Yifei She and Huangxuan Wu}, year = {2025}, eprint = {2509.00664}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2509.00664}, }