@misc{indiciaefdcec9fb0f7c, title = {Improving Speech Translation by Cross-Modal Multi-Grained Contrastive Learning}, author = {Hao Zhang and Nianwen Si and Yaqi Chen and Wenlin Zhang and Xukui Yang and Dan Qu and Wei-Qiang Zhang}, year = {2023}, doi = {10.1109/taslp.2023.3244521}, url = {https://arxiv.org/abs/2304.10309}, note = {Source identifier: 2304.10309} }