123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475 |
- WEIGHT_SHAPES = {
- "ideal": [[4 * 256 * 32, 256 * 32]],
- "mistralai/Mistral-7B-v0.1/TP1": [
- [4096, 6144],
- [4096, 4096],
- [4096, 28672],
- [14336, 4096],
- ],
- "mistralai/Mistral-7B-v0.1/TP2": [
- [4096, 3072],
- [2048, 4096],
- [4096, 14336],
- [7168, 4096],
- ],
- "mistralai/Mistral-7B-v0.1/TP4": [
- [4096, 1536],
- [1024, 4096],
- [4096, 7168],
- [3584, 4096],
- ],
- "meta-llama/Llama-2-7b-hf/TP1": [
- [4096, 12288],
- [4096, 4096],
- [4096, 22016],
- [11008, 4096],
- ],
- "meta-llama/Llama-2-7b-hf/TP2": [
- [4096, 6144],
- [2048, 4096],
- [4096, 11008],
- [5504, 4096],
- ],
- "meta-llama/Llama-2-7b-hf/TP4": [
- [4096, 3072],
- [1024, 4096],
- [4096, 5504],
- [2752, 4096],
- ],
- "meta-llama/Llama-2-13b-hf/TP1": [
- [5120, 15360],
- [5120, 5120],
- [5120, 27648],
- [13824, 5120],
- ],
- "meta-llama/Llama-2-13b-hf/TP2": [
- [5120, 7680],
- [2560, 5120],
- [5120, 13824],
- [6912, 5120],
- ],
- "meta-llama/Llama-2-13b-hf/TP4": [
- [5120, 3840],
- [1280, 5120],
- [5120, 6912],
- [3456, 5120],
- ],
- "meta-llama/Llama-2-70b-hf/TP1": [
- [8192, 10240],
- [8192, 8192],
- [8192, 57344],
- [28672, 8192],
- ],
- "meta-llama/Llama-2-70b-hf/TP2": [
- [8192, 5120],
- [4096, 8192],
- [8192, 28672],
- [14336, 8192],
- ],
- "meta-llama/Llama-2-70b-hf/TP4": [
- [8192, 2560],
- [2048, 8192],
- [8192, 14336],
- [7168, 8192],
- ],
- }
|