@misc{indiciae9bd1698e84e3, title = {TMT-VIS: Taxonomy-aware Multi-dataset Joint Training for Video Instance Segmentation}, author = {Rongkun Zheng and Lu Qi and Xi Chen and Yi Wang and Kun Wang and Yu Qiao and Hengshuang Zhao}, year = {2024}, url = {https://arxiv.org/abs/2312.06630}, note = {Source identifier: 2312.06630} }