@inproceedings{lin2026evo1,title={Evo-1: Lightweight vision-language-action model with preserved semantic alignment},author={Lin, Tao and Zhong, Y. and Du, Yuxin and Zhang, J. and Liu, J. and Chen, Y. and Gu, E. and Liu, Z. and Cai, H. and Zou, Y.},booktitle={Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition},year={2026},}
2025
VLA-Pruner: Temporal-Aware Dual-Level Visual Token Pruning for Efficient Vision-Language-Action Inference
Z. Liu, Y. Chen, H. Cai, T Lin, and 3 more authors
@article{liu2025vlapruner,title={VLA-Pruner: Temporal-Aware Dual-Level Visual Token Pruning for Efficient Vision-Language-Action Inference},author={Liu, Z. and Chen, Y. and Cai, H. and Lin, Tao and Yang, S. and Liu, Z. and Zhao, B.},journal={arXiv preprint arXiv:2511.16449},year={2025},}
Evo-0: Vision-language-action model with implicit spatial understanding
@article{lin2025evo0,title={Evo-0: Vision-language-action model with implicit spatial understanding},author={Lin, Tao and Li, Gen and Zhong, Y. and Zou, Y. and Du, Yuxin and Liu, J. and Gu, E. and Zhao, B.},journal={arXiv preprint arXiv:2507.00416},year={2025},}
Robofac: A comprehensive framework for robotic failure analysis and correction
Zewei Ye, Weifeng Lu, Minghao Ye, T Lin, and 3 more authors
@article{ye2025robofac,title={Robofac: A comprehensive framework for robotic failure analysis and correction},author={Ye, Zewei and Lu, Weifeng and Ye, Minghao and Lin, Tao and Yang, Shuo and Yan, Junchi and Zhao, Bo},journal={arXiv preprint arXiv:2505.12224},year={2025},}
Sti-bench: Are mllms ready for precise spatial-temporal world understanding?
@inproceedings{li2025stibench,title={Sti-bench: Are mllms ready for precise spatial-temporal world understanding?},author={Li, Y. and Zhang, Y. and Lin, Tao and Liu, X. R. and Cai, W. and Liu, Z. and Zhao, B.},booktitle={Proceedings of the IEEE/CVF International Conference on Computer Vision},year={2025},}