@misc{indiciae011817fa508e, title = {HORNet: Task-Guided Frame Selection for Video Question Answering with Vision-Language Models}, author = {Xiangyu Bai and Bishoy Galoaa and Sarah Ostadabbas}, year = {2026}, url = {https://arxiv.org/abs/2603.18850}, note = {Source identifier: 2603.18850} }