@misc{indiciae85eb4c9d7ce0, title = {Stacked Acoustic-and-Textual Encoding: Integrating the Pre-trained Models into Speech Translation Encoders}, author = {Chen Xu and Bojie Hu and Yanyang Li and Yuhao Zhang and shen huang and Qi Ju and Tong Xiao and Jingbo Zhu}, year = {2021}, url = {https://arxiv.org/abs/2105.05752}, note = {Source identifier: 2105.05752} }