@article{efficientvocabularyfreefinegrainedvisual, title = {Efficient Vocabulary-Free Fine-Grained Visual Recognition in the Age of Multimodal LLMs}, author = {Hari Chandana Kuchibhotla and Sai Srinivas Kancheti and Abbavaram Gowtham Reddy and Vineeth N Balasubramanian}, year = {2025}, eprint = {2505.01064}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2505.01064v1}, }