@misc{indiciaeb49665332608, title = {Language Models are Causal Knowledge Extractors for Zero-shot Video Question Answering}, author = {Hung-Ting Su and Yulei Niu and Xudong Lin and Winston H. Hsu and Shih-Fu Chang}, year = {2023}, url = {https://arxiv.org/abs/2304.03754}, note = {Source identifier: 2304.03754} }