@misc{indiciaef8d5543e91bd, title = {ViTexQA: A Multi-Frame Temporal Perception Dataset for Video Text Question Answering}, author = {Zhentao Guo and Chen Duan and Tongkun Guan and Zining Wang and Kai Zhou and Pengfei Yan}, year = {2026}, url = {https://arxiv.org/abs/2606.24602}, note = {Source identifier: 2606.24602} }