@misc{indiciae70ba86a24612, title = {Is a Video worth \$n\textbackslash{}times n\$ Images? A Highly Efficient Approach to Transformer-based Video Question Answering}, author = {Chenyang Lyu and Tianbo Ji and Yvette Graham and Jennifer Foster}, year = {2023}, url = {https://arxiv.org/abs/2305.09107}, note = {Source identifier: 2305.09107} }