@misc{indiciaefb224cdca0ac, title = {Advancing Egocentric Video Question Answering with Multimodal Large Language Models}, author = {Alkesh Patel and Vibhav Chitalia and Yinfei Yang}, year = {2025}, url = {https://arxiv.org/abs/2504.04550}, note = {Source identifier: 2504.04550} }