@misc{indiciaed953eaaabef9, title = {VideoWebArena: Evaluating Long Context Multimodal Agents with Video Understanding Web Tasks}, author = {Lawrence Jang and Yinheng Li and Dan Zhao and Charles Ding and Justin Lin and Paul Pu Liang and Rogerio Bonatti and Kazuhito Koishida}, year = {2025}, url = {https://arxiv.org/abs/2410.19100}, note = {Source identifier: 2410.19100} }