@misc{indiciae664735b2a723, title = {VideoBERT: A Joint Model for Video and Language Representation Learning}, author = {Chen Sun and Austin Myers and Carl Vondrick and Kevin Murphy and Cordelia Schmid}, year = {2019}, url = {https://arxiv.org/abs/1904.01766}, note = {Source identifier: 1904.01766} }