@misc{indiciae2dc070abe3f4, title = {Synchronizing Vision and Language: Bidirectional Token-Masking AutoEncoder for Referring Image Segmentation}, author = {Minhyeok Lee and Dogyoon Lee and Jungho Lee and Suhwan Cho and Heeseung Choi and Ig-Jae Kim and Sangyoun Lee}, year = {2023}, url = {https://arxiv.org/abs/2311.17952}, note = {Source identifier: 2311.17952} }