@misc{indiciaeb9600bf70bb1, title = {VideoOFA: Two-Stage Pre-Training for Video-to-Text Generation}, author = {Xilun Chen and Lili Yu and Wenhan Xiong and Barlas Oğuz and Yashar Mehdad and Wen-tau Yih}, year = {2023}, url = {https://arxiv.org/abs/2305.03204}, note = {Source identifier: 2305.03204} }