Download src/llama-sampler.h from OpenTransformer/llama.cpp-prismml: direct link, hf CLI and curl.
- Browser
- Download file 854 Bytes
-
https://huggingface.co/OpenTransformer/llama.cpp-prismml/resolve/main/src/llama-sampler.h
- Command line
-
hf download hf://OpenTransformer/llama.cpp-prismml/src/llama-sampler.h
-
curl -L -o llama-sampler.h https://huggingface.co/OpenTransformer/llama.cpp-prismml/resolve/main/src/llama-sampler.h
854 Bytes
| struct llama_vocab; | |
| struct llama_grammar; | |
| // sampler chain | |
| struct llama_sampler_chain { | |
| llama_sampler_chain_params params; | |
| // has .backend_init() been called? | |
| bool is_init = false; | |
| struct info { | |
| bool is_backend; | |
| llama_sampler * ptr; | |
| }; | |
| std::vector<info> samplers; | |
| // pre-allocated buffer for llama_sampler_sample to avoid repeated allocations | |
| std::vector<llama_token_data> cur; | |
| // timing | |
| mutable int64_t t_sample_us; | |
| mutable int32_t n_sample; | |
| }; | |
| struct llama_sampler * llama_sampler_init_dry_testing( | |
| int32_t context_size, | |
| float dry_multiplier, | |
| float dry_base, | |
| int32_t dry_allowed_length, | |
| int32_t dry_penalty_last_n, | |
| const std::vector<std::vector<llama_token>> & seq_breakers); | |