@misc{indiciae2a68b5358eeb, title = {Can Multimodal LLMs do Visual Temporal Understanding and Reasoning? The answer is No!}, author = {Mohamed Fazli Imam and Chenyang Lyu and Alham Fikri Aji}, year = {2025}, url = {https://arxiv.org/abs/2501.10674}, note = {Source identifier: 2501.10674} }