In the ever-evolving landscape of artificial intelligence (AI) and large language models (LLMs), handling and leveraging data effectively has become a critical challenge. Most state-of-the-art machine learning algorithms are data-centric. However, as the lifeblood of model performance, necessary data cannot always be centralized due to various factors such as privacy, regulation, geopolitics, copyright issues, and the sheer effort required to move vast datasets. In this paper, we explore how federated learning enabled by NVIDIA FLARE can address these challenges with easy and scalable integration capabilities, enabling parameter-efficient and full supervised fine-tuning of LLMs for natural language processing and biopharmaceutical applications to enhance their accuracy and robustness.
@article{arxiv.2402.07792,
title = {Empowering Federated Learning for Massive Models with NVIDIA FLARE},
author = {Holger R. Roth and Ziyue Xu and Yuan-Ting Hsieh and Adithya Renduchintala and Isaac Yang and Zhihong Zhang and Yuhong Wen and Sean Yang and Kevin Lu and Kristopher Kersten and Camir Ricketts and Daguang Xu and Chester Chen and Yan Cheng and Andrew Feng},
journal= {arXiv preprint arXiv:2402.07792},
year = {2024}
}