@misc{indiciae234fb750d11d, title = {HierarQ: Task-Aware Hierarchical Q-Former for Enhanced Video Understanding}, author = {Shehreen Azad and Vibhav Vineet and Yogesh Singh Rawat}, year = {2025}, url = {https://arxiv.org/abs/2503.08585}, note = {Source identifier: 2503.08585} }