@misc{indiciaec714e64ef702, title = {IntentVLM: Open-Vocabulary Intention Recognition through Forward-Inverse Modeling with Video-Language Models}, author = {Hamed Rahimi and Clemence Grislain and Adrien Jacquet Cretides and Olivier Sigaud and Mohamed Chetouani}, year = {2026}, url = {https://arxiv.org/abs/2604.24002}, note = {Source identifier: 2604.24002} }