@article{bilivlasceneawarevisionlanguageactionmod, title = {BiliVLA: Scene-Aware Vision-Language-Action Model with Reinforcement Learning for Autonomous Biliary Endoscopic Navigation}, author = {Jinsong Lin and Chi Kit Ng and Zhiyong Xiong and Zikang Pan and Yihan Hu and Tabassum Tamima and Ziyi Hao and Eddie Cheung and Jiewen Lai and Huxin Gao and Hongliang Ren}, year = {2026}, eprint = {2606.23531}, archivePrefix = {arXiv}, url = {https://arxiv.org/abs/2606.23531}, }