@misc{indiciaeef589d09fbc6, title = {Two Heads are Better than One: Distilling Large Language Model Features Into Small Models with Feature Decomposition and Mixture}, author = {Tianhao Fu and Xinxin Xu and Weichen Xu and Jue Chen and Ruilong Ren and Bowen Deng and Xinyu Zhao and Jian Cao and Xixin Cao}, year = {2026}, url = {https://arxiv.org/abs/2511.07110}, note = {Source identifier: 2511.07110} }