@misc{indiciae526bd695d422, title = {Voila-A: Aligning Vision-Language Models with User's Gaze Attention}, author = {Kun Yan and Lei Ji and Zeyu Wang and Yuntao Wang and Nan Duan and Shuai Ma}, year = {2023}, url = {https://arxiv.org/abs/2401.09454}, note = {Source identifier: 2401.09454} }