@misc{indiciae539e7a0662d8, title = {ActBERT: Learning Global-Local Video-Text Representations}, author = {Linchao Zhu and Yi Yang}, year = {2020}, url = {https://arxiv.org/abs/2011.07231}, note = {Source identifier: 2011.07231} }