Update README.md
Browse files
README.md
CHANGED
@@ -144,4 +144,41 @@ The following hyperparameters were used during training:
|
|
144 |
- Transformers 4.34.0
|
145 |
- Pytorch 2.0.1+cu118
|
146 |
- Datasets 2.12.0
|
147 |
-
- Tokenizers 0.14.0
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
144 |
- Transformers 4.34.0
|
145 |
- Pytorch 2.0.1+cu118
|
146 |
- Datasets 2.12.0
|
147 |
+
- Tokenizers 0.14.0
|
148 |
+
|
149 |
+
## Citation
|
150 |
+
|
151 |
+
If you find Zephyr-7B-α is useful in your work, please cite it with:
|
152 |
+
|
153 |
+
```
|
154 |
+
@misc{tunstall2023zephyr,
|
155 |
+
title={Zephyr: Direct Distillation of LM Alignment},
|
156 |
+
author={Lewis Tunstall and Edward Beeching and Nathan Lambert and Nazneen Rajani and Kashif Rasul and Younes Belkada and Shengyi Huang and Leandro von Werra and Clémentine Fourrier and Nathan Habib and Nathan Sarrazin and Omar Sanseviero and Alexander M. Rush and Thomas Wolf},
|
157 |
+
year={2023},
|
158 |
+
eprint={2310.16944},
|
159 |
+
archivePrefix={arXiv},
|
160 |
+
primaryClass={cs.LG}
|
161 |
+
}
|
162 |
+
```
|
163 |
+
|
164 |
+
If you use the UltraChat or UltraFeedback datasets, please cite the original works:
|
165 |
+
|
166 |
+
```
|
167 |
+
@misc{ding2023enhancing,
|
168 |
+
title={Enhancing Chat Language Models by Scaling High-quality Instructional Conversations},
|
169 |
+
author={Ning Ding and Yulin Chen and Bokai Xu and Yujia Qin and Zhi Zheng and Shengding Hu and Zhiyuan Liu and Maosong Sun and Bowen Zhou},
|
170 |
+
year={2023},
|
171 |
+
eprint={2305.14233},
|
172 |
+
archivePrefix={arXiv},
|
173 |
+
primaryClass={cs.CL}
|
174 |
+
}
|
175 |
+
|
176 |
+
@misc{cui2023ultrafeedback,
|
177 |
+
title={UltraFeedback: Boosting Language Models with High-quality Feedback},
|
178 |
+
author={Ganqu Cui and Lifan Yuan and Ning Ding and Guanming Yao and Wei Zhu and Yuan Ni and Guotong Xie and Zhiyuan Liu and Maosong Sun},
|
179 |
+
year={2023},
|
180 |
+
eprint={2310.01377},
|
181 |
+
archivePrefix={arXiv},
|
182 |
+
primaryClass={cs.CL}
|
183 |
+
}
|
184 |
+
```
|