About Me

Hang Zhang

I am Hang Zhang (张航), Senior Principal Engineer and Senior Director at XPeng, and Head of the VLA Foundation Model Team. We build XPeng VLA 2.0, the pure-vision, end-to-end model behind XPeng's production driver-assistance system in 300+ Chinese cities, overseas vehicles, and Robotaxi. I drove the program from early research to large-scale production, reaching an industry-leading position among domestic peers. Previously at Cruise, Meta, and Amazon AI, where we developed ResNeSt. Named a Top 50 Intelligent Driving Developer in China (2025). My work has been cited 12,000+ times on Google Scholar, with 10,000+ GitHub stars.

XPeng · VLA Foundation Model ex-Cruise ex-Meta ex-Amazon AI PhD, Rutgers

News


Selected Publications

CLAP

CLAP: Unsupervised 3D Representation Learning for Fusion 3D Perception via Curvature Sampling and Prototype Learning
Runjian Chen, Hang Zhang, Avinash Ravichandran, Hyoungseob Park, Wenqi Shao, Alex Wong, Ping Luo
International Conference on Learning Representations (ICLR), 2026

paper abstract bibtex

@inproceedings{chen2026clap,
  title={CLAP: Unsupervised 3D Representation Learning for Fusion 3D Perception via Curvature Sampling and Prototype Learning},
  author={Chen, Runjian and Zhang, Hang and Ravichandran, Avinash and Park, Hyoungseob and Shao, Wenqi and Wong, Alex and Luo, Ping},
  booktitle={International Conference on Learning Representations (ICLR)},
  year={2026}
}
    

TuringViT

TuringViT: Making SOTA Vision Transformers Accessible to All
Foundation Model Team, XPeng Inc.
arXiv, 2026

paper abstract bibtex

@article{wu2026turingvit,
  title={TuringViT: Making SOTA Vision Transformers Accessible to All},
  author={Wu, Qiman and Chen, Hanlin and Chen, Lyujie and Xin, Rui and Zheng, Jianlei and Wang, Mingyuan and Hu, Jiahui and Zhu, Da and Ma, Yuecheng and Wei, Yuhua and Wang, Yizhao and Zhou, Hua and Zhang, Yuheng and Liu, Anhua and Tang, Shaman and He, Yue and Diao, Pengfei and Su, Shuang and Xin, Haotong and Huang, Weichao and Zhang, Hang and Liu, Xianming},
  journal={arXiv preprint arXiv:2606.24253},
  year={2026}
}
    

X-Mind

X-Mind: Efficient Visual Chain-of-Thought via Predictive World Model for End-to-End Driving
Foundation Model Team, XPeng Inc.
arXiv, 2026

paper abstract bibtex

@article{zhao2026xmind,
  title={X-Mind: Efficient Visual Chain-of-Thought via Predictive World Model for End-to-End Driving},
  author={Zhao, Bohao and Wei, Chengrui and Jiang, Guangfeng and Liu, Ruixin and Lv, Xuejie and Liang, Liu and Deng, Sutao and Fan, Xiuyang and Zheng, Pengkun and Zhou, Jinyun and Guo, Rui and Liu, Hanpeng and Zheng, Yutong and Guo, Yi and Zheng, Xinlong and Luo, Qingyu and Ding, Zhuangzhuang and Zhang, Yu and Zhang, Hang and Liu, Xianming},
  journal={arXiv preprint arXiv:2606.28758},
  year={2026}
}
    

li21

Neural Architecture Search for Multiple Tasks in One Run
Bichen Wu, Chaojian Li, Hang Zhang, Xiaoliang Dai, Matthew Yu, Jialiang Wang, Yingyan Lin, Peter Vajda
arXiv, 11/2021

paper abstract bibtex

@article{wu2021fbnetv5,
  title={FBNetV5: Neural Architecture Search for Multiple Tasks in One Run},
  author={Wu, Bichen and Li, Chaojian and Zhang, Hang and Dai, Xiaoliang and Zhang, Peizhao and Yu, Matthew and Wang, Jialiang and Lin, Yingyan and Vajda, Peter},
  journal={arXiv preprint arXiv:2111.10007},
  year={2021}
}
    

zhang20

ResNeSt: Split-Attention Networks
Hang Zhang, Chongruo Wu, Zhongyue Zhang, Yi Zhu, Haibin Lin, Zhi Zhang, Yue Sun, Tong He, Jonas Mueller, R. Manmatha, Mu Li, Alex Smola
IEEE Conference on Computer Vision and Pattern Recognition Workshops (CVPRW), 2022
arXiv, 04/2020

paper abstract bibtex slides code

@InProceedings{Zhang_2022_CVPR,
    author    = {Zhang, Hang and Wu, Chongruo and Zhang, Zhongyue and Zhu, Yi and Lin, Haibin and Zhang, Zhi and Sun, Yue and He, Tong and Mueller, Jonas and Manmatha, R. and Li, Mu and Smola, Alexander},
    title     = {ResNeSt: Split-Attention Networks},
    booktitle = {Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition (CVPR) Workshops},
    month     = {June},
    year      = {2022},
    pages     = {2736-2746}
}
    

zhang19

Co-occurrent Features in Semantic Segmentation
Hang Zhang, Han Zhang, Chenguang Wang, Junyuan Xie
IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2019

paper abstract bibtex

@InProceedings{Zhang_2019_CVPR,
author = {Hang Zhang and Han Zhang and Chenguang Wang and Junyuan Xie},
title = {Co-occurrent Features in Semantic Segmentation},
booktitle = {The IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
year = {2019}
}
    

xie19

Bag of Tricks for Image Classification with Convolutional Neural Networks
Tong He, Zhi Zhang, Hang Zhang, Zhongyue Zhang, Junyuan Xie, Mu Li
IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2019

paper abstract bibtex code

@InProceedings{Xie2018bags,
  title={Bag of Tricks to Train Convolutional Neural Networks for Image Classification},
  author={Tong He and Zhi Zhang and Hang Zhang and Zhongyue Zhang and Junyuan Xie and Mu Li},
  booktitle = {The IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
  year={2018}
}
    

zhang18

Context Encoding for Semantic Segmentation
Hang Zhang, Kristin Dana, Jianping Shi, Zhongyue Zhang, Xiaogang Wang, Ambrish Tyagi, Amit Agrawal
IEEE Conference on Computer Vision and Pattern Recognition
(CVPR)
, 2018 Oral (70/3309=2.1%)

paper abstract bibtex code talk slides

@InProceedings{Zhang_2018_CVPR,
author = {Zhang, Hang and Dana, Kristin and Shi, Jianping and Zhang, Zhongyue and Wang, Xiaogang and Tyagi, Ambrish and Agrawal, Amit},
title = {Context Encoding for Semantic Segmentation},
booktitle = {The IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
month = {June},
year = {2018}
}
    

zhang17

Multi-style Generative Network for Real-time Transfer
Hang Zhang, Kristin Dana
European Conference on Computer Vision Workshops (ECCVW), 2018
arXiv, 03/2017

paper abstract bibtex code video project poster

@article{zhang2017multistyle,
title={Multi-style Generative Network for Real-time Transfer},
author={Zhang, Hang and Dana, Kristin},
journal={arXiv preprint arXiv:1703.06953},
year={2017}
}
    

zhang17

Deep TEN: Texture Encoding Network
Hang Zhang, Jia Xue, Kristin Dana
IEEE Conference on Computer Vision and Pattern Recognition (CVPR), 2017

paper abstract bibtex code blog poster slides

@InProceedings{Zhang_2017_CVPR,
author = {Zhang, Hang and Xue, Jia and Dana, Kristin},
title = {Deep TEN: Texture Encoding Network},
booktitle = {The IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
month = {July},
year = {2017}
}
    


sitemap  ·  sitemap.xml