@article{howwellcangeneralvisionlanguagemodels, title = {How Well Can General Vision-Language Models Learn Medicine By Watching Public Educational Videos?}, author = {Rahul Thapa and Andrew Li and Qingyang Wu and Bryan He and Yuki Sahashi and Christina Binder and Angela Zhang and Ben Athiwaratkun and Shuaiwen Leon Song and David Ouyang and James Zou}, year = {2025}, eprint = {2504.14391}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2504.14391v1}, }