server : enable -td and -tbd parameters (#15172)

This commit is contained in:
Sigbjørn Skjæret
2025-08-13 15:43:00 +02:00
committed by GitHub
parent c24f4e2688
commit b3e16665e1
2 changed files with 4 additions and 2 deletions

View File

@@ -2015,6 +2015,8 @@ struct server_context {
params_dft.cache_type_k = params_base.speculative.cache_type_k;
params_dft.cache_type_v = params_base.speculative.cache_type_v;
params_dft.cpuparams.n_threads = params_base.speculative.cpuparams.n_threads;
params_dft.cpuparams_batch.n_threads = params_base.speculative.cpuparams_batch.n_threads;
params_dft.tensor_buft_overrides = params_base.speculative.tensor_buft_overrides;
llama_init_dft = common_init_from_params(params_dft);