@article{humanvlmfoundationforhumanscenevision, title = {HumanVLM: Foundation for Human-Scene Vision-Language Model}, author = {Dawei Dai and Xu Long and Li Yutang and Zhang YuanHui and Shuyin Xia}, year = {2024}, eprint = {2411.03034}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2411.03034v1}, }