@misc{indiciae3f9b8a94629f, title = {Unified Embedding Alignment for Open-Vocabulary Video Instance Segmentation}, author = {Hao Fang and Peng Wu and Yawei Li and Xinxin Zhang and Xiankai Lu}, year = {2024}, url = {https://arxiv.org/abs/2407.07427}, note = {Source identifier: 2407.07427} }