LM/hf: Difference between revisions
< LM
Jump to navigation
Jump to search
| (9 intermediate revisions by the same user not shown) | |||
| Line 17: | Line 17: | ||
</syntaxhighlight> | </syntaxhighlight> | ||
|- | |- | ||
| search || | | search for vLLM || | ||
<syntaxhighlight lang="bash"> | <syntaxhighlight lang="bash"> | ||
# search by population | # search by population | ||
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --limit | hf models ls --search "Qwen3.8 27B MixedINT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10 | ||
# search by population | |||
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10 | |||
# search for latest publish | |||
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort created_at --no-truncate --no-gated --limit 10 | |||
# search for latest tunning | |||
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort last_modified --no-truncate --no-gated --limit 10 | |||
</syntaxhighlight> | |||
|- | |||
| search for llama.cpp || | |||
<syntaxhighlight lang="bash"> | |||
# search by population | |||
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10 | |||
# search for latest publish | # search for latest publish | ||
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort created_at --no-truncate --limit | hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort created_at --no-truncate --no-gated --limit 10 | ||
# search for latest tunning | # search for latest tunning | ||
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort last_modified --no-truncate --limit | hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort last_modified --no-truncate --no-gated --limit 10 | ||
</syntaxhighlight> | |||
|- | |||
| search for OpenVINO || | |||
<syntaxhighlight lang="bash"> | |||
# search OpenVINO format | # search OpenVINO format | ||
hf models ls --search "Qwen3.8" --filter openvino --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --limit | hf models ls --search "Qwen3.8" --filter openvino --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10 | ||
</syntaxhighlight> | </syntaxhighlight> | ||
|- | |- | ||
Latest revision as of 07:17, 24 September 2026
Quick references
| Purpose | Command |
|---|---|
| cache management |
hf cache list
hf cache rm <model id>
hf cache prune
|
| fix WiFi problem |
HF_XET_FIXED_DOWNLOAD_CONCURRENCY=10 hf download "unsloth/Qwen3-Coder-30B-A3B-Instruct-GGUF" --include "*UD-Q4_K_XL*"
HF_XET_FIXED_DOWNLOAD_CONCURRENCY=10 hf download "unsloth/Devstral-Small-2-24B-Instruct-2512-GGUF" --include "*UD-Q4_K_XL*"
|
| search for vLLM |
# search by population
hf models ls --search "Qwen3.8 27B MixedINT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10
# search by population
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10
# search for latest publish
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort created_at --no-truncate --no-gated --limit 10
# search for latest tunning
hf models ls --search "Qwen3.8 27B INT4 AutoRound" --apps vllm --expand "downloads,likes,createdAt,lastModified" --sort last_modified --no-truncate --no-gated --limit 10
|
| search for llama.cpp |
# search by population
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10
# search for latest publish
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort created_at --no-truncate --no-gated --limit 10
# search for latest tunning
hf models ls --search "Muse-Glimmer" --apps llama.cpp --expand "downloads,likes,createdAt,lastModified" --sort last_modified --no-truncate --no-gated --limit 10
|
| search for OpenVINO |
# search OpenVINO format
hf models ls --search "Qwen3.8" --filter openvino --expand "downloads,likes,createdAt,lastModified" --sort downloads --no-truncate --no-gated --limit 10
|
| optimize Muse Glimmer |
hf models ls -h unsloth/Muse-Glimmer-30B-GGUF
hf download unsloth/Muse-Glimmer-30B-GGUF --include "*UD-Q4_K_XL*"
hf download unsloth/Muse-Glimmer-30B-GGUF --include "*UD-Q5_K_XL*"
hf download unsloth/Muse-Glimmer-30B-GGUF --include "*UD-Q6_K_XL*"
hf download unsloth/Muse-Glimmer-30B-GGUF --include "dflash-kquant.gguf"
hf download unsloth/Muse-Glimmer-30B-GGUF --include "mmproj-kquant.gguf"
|
tree ~/.cache/huggingface/hub/models--unsloth--Devstral-Small-2-24B-Instruct-2512-GGUF/snapshots
tree ~/.cache/huggingface/hub/models--unsloth--Qwen3-Coder-30B-A3B-Instruct-GGUF/snapshots
|
Install
sudo apt update
sudo apt install pipx -y
pipx ensurepath
pipx install "huggingface_hub[cli]"