@misc{indiciae85f32b6b4237, title = {Attention-driven GUI Grounding: Leveraging Pretrained Multimodal Large Language Models without Fine-Tuning}, author = {Hai-Ming Xu and Qi Chen and Lei Wang and Lingqiao Liu}, year = {2024}, url = {https://arxiv.org/abs/2412.10840}, note = {Source identifier: 2412.10840} }