@misc{indiciae707ee7538e4d, title = {Video-Text Representation Learning via Differentiable Weak Temporal Alignment}, author = {Dohwan Ko and Joonmyung Choi and Juyeon Ko and Shinyeong Noh and Kyoung-Woon On and Eun-Sol Kim and Hyunwoo J. Kim}, year = {2022}, url = {https://arxiv.org/abs/2203.16784}, note = {Source identifier: 2203.16784} }