2024-02-19 23:14:04 +01:00
|
|
|
accelerate==0.27.*
|
2024-04-05 18:26:49 +02:00
|
|
|
aqlm[gpu,cpu]==1.1.3; platform_system == "Linux"
|
2024-03-10 16:30:53 +01:00
|
|
|
bitsandbytes==0.43.*
|
2023-04-17 04:26:52 +02:00
|
|
|
colorama
|
2023-04-03 01:34:25 +02:00
|
|
|
datasets
|
2023-05-29 15:20:18 +02:00
|
|
|
einops
|
2024-04-11 07:24:53 +02:00
|
|
|
gradio==4.26.*
|
2024-03-05 11:33:51 +01:00
|
|
|
hqq==0.1.5
|
2024-01-14 01:47:13 +01:00
|
|
|
jinja2==3.1.2
|
2023-12-24 18:22:31 +01:00
|
|
|
lm_eval==0.3.0
|
2023-03-15 16:40:03 +01:00
|
|
|
markdown
|
2024-03-11 00:13:29 +01:00
|
|
|
numba==0.59.*
|
2024-02-13 20:26:35 +01:00
|
|
|
numpy==1.26.*
|
2024-02-19 23:15:21 +01:00
|
|
|
optimum==1.17.*
|
2023-04-21 05:20:33 +02:00
|
|
|
pandas
|
2024-03-05 11:56:37 +01:00
|
|
|
peft==0.8.*
|
2023-04-08 23:48:46 +02:00
|
|
|
Pillow>=9.5.0
|
2024-04-07 04:02:20 +02:00
|
|
|
psutil
|
2023-04-21 05:20:33 +02:00
|
|
|
pyyaml
|
2023-03-11 18:47:30 +01:00
|
|
|
requests
|
2023-12-20 06:58:36 +01:00
|
|
|
rich
|
2024-02-05 03:40:25 +01:00
|
|
|
safetensors==0.4.*
|
2023-05-26 04:26:25 +02:00
|
|
|
scipy
|
2023-07-20 04:31:19 +02:00
|
|
|
sentencepiece
|
2023-07-12 16:53:31 +02:00
|
|
|
tensorboard
|
2024-04-19 00:55:34 +02:00
|
|
|
transformers==4.40.*
|
2023-07-20 04:31:19 +02:00
|
|
|
tqdm
|
|
|
|
wandb
|
2023-08-09 17:07:55 +02:00
|
|
|
|
2024-03-04 08:46:39 +01:00
|
|
|
# API
|
|
|
|
SpeechRecognition==3.10.0
|
|
|
|
flask_cloudflared==0.0.14
|
2024-04-18 20:05:00 +02:00
|
|
|
sse-starlette==1.6.5
|
2024-03-04 08:46:39 +01:00
|
|
|
tiktoken
|
2024-05-23 14:09:54 +02:00
|
|
|
soundfile
|
|
|
|
openai-whisper
|
2024-03-04 08:46:39 +01:00
|
|
|
|
2024-04-30 14:11:31 +02:00
|
|
|
# llama-cpp-python (CPU only, AVX2)
|
2024-05-11 19:53:19 +02:00
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/cpu/llama_cpp_python-0.2.73+cpuavx2-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/cpu/llama_cpp_python-0.2.73+cpuavx2-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/cpu/llama_cpp_python-0.2.73+cpuavx2-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/cpu/llama_cpp_python-0.2.73+cpuavx2-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
2024-04-30 14:11:31 +02:00
|
|
|
|
|
|
|
# llama-cpp-python (CUDA, no tensor cores)
|
2024-05-11 19:53:19 +02:00
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda-0.2.73+cu121-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda-0.2.73+cu121-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda-0.2.73+cu121-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda-0.2.73+cu121-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
2024-04-30 14:11:31 +02:00
|
|
|
|
|
|
|
# llama-cpp-python (CUDA, tensor cores)
|
2024-05-11 19:53:19 +02:00
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda_tensorcores-0.2.73+cu121-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda_tensorcores-0.2.73+cu121-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda_tensorcores-0.2.73+cu121-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/llama-cpp-python-cuBLAS-wheels/releases/download/textgen-webui/llama_cpp_python_cuda_tensorcores-0.2.73+cu121-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
2023-12-19 21:30:53 +01:00
|
|
|
|
2023-09-24 14:58:29 +02:00
|
|
|
# CUDA wheels
|
2023-12-15 15:18:49 +01:00
|
|
|
https://github.com/jllllll/AutoGPTQ/releases/download/v0.6.0/auto_gptq-0.6.0+cu121-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/jllllll/AutoGPTQ/releases/download/v0.6.0/auto_gptq-0.6.0+cu121-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/jllllll/AutoGPTQ/releases/download/v0.6.0/auto_gptq-0.6.0+cu121-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/jllllll/AutoGPTQ/releases/download/v0.6.0/auto_gptq-0.6.0+cu121-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
2024-05-02 01:20:50 +02:00
|
|
|
https://github.com/oobabooga/exllamav2/releases/download/v0.0.20/exllamav2-0.0.20+cu121-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/exllamav2/releases/download/v0.0.20/exllamav2-0.0.20+cu121-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/oobabooga/exllamav2/releases/download/v0.0.20/exllamav2-0.0.20+cu121-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/exllamav2/releases/download/v0.0.20/exllamav2-0.0.20+cu121-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
|
|
|
https://github.com/oobabooga/exllamav2/releases/download/v0.0.20/exllamav2-0.0.20-py3-none-any.whl; platform_system == "Linux" and platform_machine != "x86_64"
|
2024-03-03 23:40:32 +01:00
|
|
|
https://github.com/oobabooga/flash-attention/releases/download/v2.5.6/flash_attn-2.5.6+cu122torch2.2.0cxx11abiFALSE-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/oobabooga/flash-attention/releases/download/v2.5.6/flash_attn-2.5.6+cu122torch2.2.0cxx11abiFALSE-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/Dao-AILab/flash-attention/releases/download/v2.5.6/flash_attn-2.5.6+cu122torch2.2cxx11abiFALSE-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/Dao-AILab/flash-attention/releases/download/v2.5.6/flash_attn-2.5.6+cu122torch2.2cxx11abiFALSE-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
2023-10-21 08:46:23 +02:00
|
|
|
https://github.com/jllllll/GPTQ-for-LLaMa-CUDA/releases/download/0.1.1/gptq_for_llama-0.1.1+cu121-cp311-cp311-win_amd64.whl; platform_system == "Windows" and python_version == "3.11"
|
|
|
|
https://github.com/jllllll/GPTQ-for-LLaMa-CUDA/releases/download/0.1.1/gptq_for_llama-0.1.1+cu121-cp310-cp310-win_amd64.whl; platform_system == "Windows" and python_version == "3.10"
|
|
|
|
https://github.com/jllllll/GPTQ-for-LLaMa-CUDA/releases/download/0.1.1/gptq_for_llama-0.1.1+cu121-cp311-cp311-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.11"
|
|
|
|
https://github.com/jllllll/GPTQ-for-LLaMa-CUDA/releases/download/0.1.1/gptq_for_llama-0.1.1+cu121-cp310-cp310-linux_x86_64.whl; platform_system == "Linux" and platform_machine == "x86_64" and python_version == "3.10"
|
2024-03-08 23:54:56 +01:00
|
|
|
autoawq==0.2.3; platform_system == "Linux" or platform_system == "Windows"
|