流程节点完善
This commit is contained in:
@@ -120,7 +120,6 @@ print_info: EOG token = 151663 '<|repo_name|>'
|
||||
print_info: EOG token = 151664 '<|file_sep|>'
|
||||
print_info: max token length = 256
|
||||
load_tensors: loading model tensors, this can take a while... (mmap = true)
|
||||
srv log_server_r: request: GET /health 127.0.0.1 503
|
||||
load_tensors: offloading 36 repeating layers to GPU
|
||||
load_tensors: offloaded 36/37 layers to GPU
|
||||
load_tensors: CUDA0 model buffer size = 2445.68 MiB
|
||||
@@ -221,227 +220,37 @@ How are you?<|im_end|>
|
||||
'
|
||||
main: server is listening on http://0.0.0.0:8090 - starting the main loop
|
||||
srv update_slots: all slots are idle
|
||||
slot launch_slot_: id 0 | task 0 | processing task
|
||||
slot update_slots: id 0 | task 0 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 38
|
||||
slot update_slots: id 0 | task 0 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 0 | prompt processing progress, n_past = 38, n_tokens = 38, progress = 1.000000
|
||||
slot update_slots: id 0 | task 0 | prompt done, n_past = 38, n_tokens = 38
|
||||
slot release: id 0 | task 0 | stop processing: n_past = 38, truncated = 0
|
||||
srv log_server_r: request: GET /health 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 1 | processing task
|
||||
slot update_slots: id 0 | task 1 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 50
|
||||
slot update_slots: id 0 | task 1 | kv cache rm [4, end)
|
||||
slot update_slots: id 0 | task 1 | prompt processing progress, n_past = 50, n_tokens = 46, progress = 0.920000
|
||||
slot update_slots: id 0 | task 1 | prompt done, n_past = 50, n_tokens = 46
|
||||
slot release: id 0 | task 1 | stop processing: n_past = 50, truncated = 0
|
||||
slot launch_slot_: id 0 | task 2 | processing task
|
||||
slot update_slots: id 0 | task 2 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 142
|
||||
slot update_slots: id 0 | task 2 | kv cache rm [4, end)
|
||||
slot update_slots: id 0 | task 2 | prompt processing progress, n_past = 142, n_tokens = 138, progress = 0.971831
|
||||
slot update_slots: id 0 | task 2 | prompt done, n_past = 142, n_tokens = 138
|
||||
slot release: id 0 | task 2 | stop processing: n_past = 142, truncated = 0
|
||||
slot launch_slot_: id 0 | task 3 | processing task
|
||||
slot update_slots: id 0 | task 3 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 27
|
||||
slot update_slots: id 0 | task 3 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 3 | prompt processing progress, n_past = 27, n_tokens = 27, progress = 1.000000
|
||||
slot update_slots: id 0 | task 3 | prompt done, n_past = 27, n_tokens = 27
|
||||
slot release: id 0 | task 3 | stop processing: n_past = 27, truncated = 0
|
||||
slot update_slots: id 0 | task 1 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 1 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 1 | prompt processing progress, n_past = 9, n_tokens = 9, progress = 1.000000
|
||||
slot update_slots: id 0 | task 1 | prompt done, n_past = 9, n_tokens = 9
|
||||
slot release: id 0 | task 1 | stop processing: n_past = 9, truncated = 0
|
||||
slot launch_slot_: id 0 | task 0 | processing task
|
||||
slot update_slots: id 0 | task 0 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 0 | need to evaluate at least 1 token for each active slot, n_past = 9, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 0 | kv cache rm [8, end)
|
||||
slot update_slots: id 0 | task 0 | prompt processing progress, n_past = 9, n_tokens = 1, progress = 0.111111
|
||||
slot update_slots: id 0 | task 0 | prompt done, n_past = 9, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 0 | stop processing: n_past = 9, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 4 | processing task
|
||||
slot update_slots: id 0 | task 4 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 20
|
||||
slot update_slots: id 0 | task 4 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 4 | prompt processing progress, n_past = 20, n_tokens = 20, progress = 1.000000
|
||||
slot update_slots: id 0 | task 4 | prompt done, n_past = 20, n_tokens = 20
|
||||
slot release: id 0 | task 4 | stop processing: n_past = 20, truncated = 0
|
||||
slot launch_slot_: id 0 | task 5 | processing task
|
||||
slot update_slots: id 0 | task 5 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 5 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 5 | prompt processing progress, n_past = 17, n_tokens = 17, progress = 1.000000
|
||||
slot update_slots: id 0 | task 5 | prompt done, n_past = 17, n_tokens = 17
|
||||
slot release: id 0 | task 5 | stop processing: n_past = 17, truncated = 0
|
||||
slot update_slots: id 0 | task 4 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 4 | need to evaluate at least 1 token for each active slot, n_past = 9, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 4 | kv cache rm [8, end)
|
||||
slot update_slots: id 0 | task 4 | prompt processing progress, n_past = 9, n_tokens = 1, progress = 0.111111
|
||||
slot update_slots: id 0 | task 4 | prompt done, n_past = 9, n_tokens = 1
|
||||
slot release: id 0 | task 4 | stop processing: n_past = 9, truncated = 0
|
||||
slot launch_slot_: id 0 | task 6 | processing task
|
||||
slot update_slots: id 0 | task 6 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 10
|
||||
slot update_slots: id 0 | task 6 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 6 | prompt processing progress, n_past = 10, n_tokens = 10, progress = 1.000000
|
||||
slot update_slots: id 0 | task 6 | prompt done, n_past = 10, n_tokens = 10
|
||||
slot release: id 0 | task 6 | stop processing: n_past = 10, truncated = 0
|
||||
slot launch_slot_: id 0 | task 7 | processing task
|
||||
slot update_slots: id 0 | task 7 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 7 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 7 | prompt processing progress, n_past = 17, n_tokens = 17, progress = 1.000000
|
||||
slot update_slots: id 0 | task 7 | prompt done, n_past = 17, n_tokens = 17
|
||||
slot release: id 0 | task 7 | stop processing: n_past = 17, truncated = 0
|
||||
slot launch_slot_: id 0 | task 8 | processing task
|
||||
slot update_slots: id 0 | task 8 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 18
|
||||
slot update_slots: id 0 | task 8 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 8 | prompt processing progress, n_past = 18, n_tokens = 18, progress = 1.000000
|
||||
slot update_slots: id 0 | task 8 | prompt done, n_past = 18, n_tokens = 18
|
||||
slot release: id 0 | task 8 | stop processing: n_past = 18, truncated = 0
|
||||
slot launch_slot_: id 0 | task 9 | processing task
|
||||
slot update_slots: id 0 | task 9 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 23
|
||||
slot update_slots: id 0 | task 9 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 9 | prompt processing progress, n_past = 23, n_tokens = 23, progress = 1.000000
|
||||
slot update_slots: id 0 | task 9 | prompt done, n_past = 23, n_tokens = 23
|
||||
slot release: id 0 | task 9 | stop processing: n_past = 23, truncated = 0
|
||||
slot launch_slot_: id 0 | task 10 | processing task
|
||||
slot update_slots: id 0 | task 10 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 25
|
||||
slot update_slots: id 0 | task 10 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 10 | prompt processing progress, n_past = 25, n_tokens = 25, progress = 1.000000
|
||||
slot update_slots: id 0 | task 10 | prompt done, n_past = 25, n_tokens = 25
|
||||
slot release: id 0 | task 10 | stop processing: n_past = 25, truncated = 0
|
||||
slot launch_slot_: id 0 | task 11 | processing task
|
||||
slot update_slots: id 0 | task 11 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 24
|
||||
slot update_slots: id 0 | task 11 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 11 | prompt processing progress, n_past = 24, n_tokens = 24, progress = 1.000000
|
||||
slot update_slots: id 0 | task 11 | prompt done, n_past = 24, n_tokens = 24
|
||||
slot release: id 0 | task 11 | stop processing: n_past = 24, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 24 | processing task
|
||||
slot update_slots: id 0 | task 24 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 38
|
||||
slot update_slots: id 0 | task 24 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 24 | prompt processing progress, n_past = 38, n_tokens = 38, progress = 1.000000
|
||||
slot update_slots: id 0 | task 24 | prompt done, n_past = 38, n_tokens = 38
|
||||
slot release: id 0 | task 24 | stop processing: n_past = 38, truncated = 0
|
||||
slot launch_slot_: id 0 | task 25 | processing task
|
||||
slot update_slots: id 0 | task 25 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 50
|
||||
slot update_slots: id 0 | task 25 | kv cache rm [4, end)
|
||||
slot update_slots: id 0 | task 25 | prompt processing progress, n_past = 50, n_tokens = 46, progress = 0.920000
|
||||
slot update_slots: id 0 | task 25 | prompt done, n_past = 50, n_tokens = 46
|
||||
slot release: id 0 | task 25 | stop processing: n_past = 50, truncated = 0
|
||||
slot launch_slot_: id 0 | task 26 | processing task
|
||||
slot update_slots: id 0 | task 26 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 142
|
||||
slot update_slots: id 0 | task 26 | kv cache rm [4, end)
|
||||
slot update_slots: id 0 | task 26 | prompt processing progress, n_past = 142, n_tokens = 138, progress = 0.971831
|
||||
slot update_slots: id 0 | task 26 | prompt done, n_past = 142, n_tokens = 138
|
||||
slot release: id 0 | task 26 | stop processing: n_past = 142, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 30 | processing task
|
||||
slot update_slots: id 0 | task 30 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 27
|
||||
slot update_slots: id 0 | task 30 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 30 | prompt processing progress, n_past = 27, n_tokens = 27, progress = 1.000000
|
||||
slot update_slots: id 0 | task 30 | prompt done, n_past = 27, n_tokens = 27
|
||||
slot release: id 0 | task 30 | stop processing: n_past = 27, truncated = 0
|
||||
slot launch_slot_: id 0 | task 31 | processing task
|
||||
slot update_slots: id 0 | task 31 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 20
|
||||
slot update_slots: id 0 | task 31 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 31 | prompt processing progress, n_past = 20, n_tokens = 20, progress = 1.000000
|
||||
slot update_slots: id 0 | task 31 | prompt done, n_past = 20, n_tokens = 20
|
||||
slot release: id 0 | task 31 | stop processing: n_past = 20, truncated = 0
|
||||
slot launch_slot_: id 0 | task 32 | processing task
|
||||
slot update_slots: id 0 | task 32 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 32 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 32 | prompt processing progress, n_past = 17, n_tokens = 17, progress = 1.000000
|
||||
slot update_slots: id 0 | task 32 | prompt done, n_past = 17, n_tokens = 17
|
||||
slot release: id 0 | task 32 | stop processing: n_past = 17, truncated = 0
|
||||
slot launch_slot_: id 0 | task 33 | processing task
|
||||
slot update_slots: id 0 | task 33 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 10
|
||||
slot update_slots: id 0 | task 33 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 33 | prompt processing progress, n_past = 10, n_tokens = 10, progress = 1.000000
|
||||
slot update_slots: id 0 | task 33 | prompt done, n_past = 10, n_tokens = 10
|
||||
slot release: id 0 | task 33 | stop processing: n_past = 10, truncated = 0
|
||||
slot launch_slot_: id 0 | task 34 | processing task
|
||||
slot update_slots: id 0 | task 34 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 34 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 34 | prompt processing progress, n_past = 17, n_tokens = 17, progress = 1.000000
|
||||
slot update_slots: id 0 | task 34 | prompt done, n_past = 17, n_tokens = 17
|
||||
slot release: id 0 | task 34 | stop processing: n_past = 17, truncated = 0
|
||||
slot launch_slot_: id 0 | task 35 | processing task
|
||||
slot update_slots: id 0 | task 35 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 18
|
||||
slot update_slots: id 0 | task 35 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 35 | prompt processing progress, n_past = 18, n_tokens = 18, progress = 1.000000
|
||||
slot update_slots: id 0 | task 35 | prompt done, n_past = 18, n_tokens = 18
|
||||
slot release: id 0 | task 35 | stop processing: n_past = 18, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 42 | processing task
|
||||
slot update_slots: id 0 | task 42 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 23
|
||||
slot update_slots: id 0 | task 42 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 42 | prompt processing progress, n_past = 23, n_tokens = 23, progress = 1.000000
|
||||
slot update_slots: id 0 | task 42 | prompt done, n_past = 23, n_tokens = 23
|
||||
slot release: id 0 | task 42 | stop processing: n_past = 23, truncated = 0
|
||||
slot launch_slot_: id 0 | task 43 | processing task
|
||||
slot update_slots: id 0 | task 43 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 25
|
||||
slot update_slots: id 0 | task 43 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 43 | prompt processing progress, n_past = 25, n_tokens = 25, progress = 1.000000
|
||||
slot update_slots: id 0 | task 43 | prompt done, n_past = 25, n_tokens = 25
|
||||
slot release: id 0 | task 43 | stop processing: n_past = 25, truncated = 0
|
||||
slot launch_slot_: id 0 | task 44 | processing task
|
||||
slot update_slots: id 0 | task 44 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 24
|
||||
slot update_slots: id 0 | task 44 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 44 | prompt processing progress, n_past = 24, n_tokens = 24, progress = 1.000000
|
||||
slot update_slots: id 0 | task 44 | prompt done, n_past = 24, n_tokens = 24
|
||||
slot release: id 0 | task 44 | stop processing: n_past = 24, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 48 | processing task
|
||||
slot update_slots: id 0 | task 48 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 2
|
||||
slot update_slots: id 0 | task 48 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 48 | prompt processing progress, n_past = 2, n_tokens = 2, progress = 1.000000
|
||||
slot update_slots: id 0 | task 48 | prompt done, n_past = 2, n_tokens = 2
|
||||
slot release: id 0 | task 48 | stop processing: n_past = 2, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 50 | processing task
|
||||
slot update_slots: id 0 | task 50 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 12
|
||||
slot update_slots: id 0 | task 50 | kv cache rm [0, end)
|
||||
slot update_slots: id 0 | task 50 | prompt processing progress, n_past = 12, n_tokens = 12, progress = 1.000000
|
||||
slot update_slots: id 0 | task 50 | prompt done, n_past = 12, n_tokens = 12
|
||||
slot release: id 0 | task 50 | stop processing: n_past = 12, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 52 | processing task
|
||||
slot update_slots: id 0 | task 52 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 52 | kv cache rm [3, end)
|
||||
slot update_slots: id 0 | task 52 | prompt processing progress, n_past = 17, n_tokens = 14, progress = 0.823529
|
||||
slot update_slots: id 0 | task 52 | prompt done, n_past = 17, n_tokens = 14
|
||||
slot release: id 0 | task 52 | stop processing: n_past = 17, truncated = 0
|
||||
slot launch_slot_: id 0 | task 53 | processing task
|
||||
slot update_slots: id 0 | task 53 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 53 | need to evaluate at least 1 token for each active slot, n_past = 17, n_prompt_tokens = 17
|
||||
slot update_slots: id 0 | task 53 | kv cache rm [16, end)
|
||||
slot update_slots: id 0 | task 53 | prompt processing progress, n_past = 17, n_tokens = 1, progress = 0.058824
|
||||
slot update_slots: id 0 | task 53 | prompt done, n_past = 17, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 53 | stop processing: n_past = 17, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 56 | processing task
|
||||
slot update_slots: id 0 | task 56 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 16
|
||||
slot update_slots: id 0 | task 56 | kv cache rm [5, end)
|
||||
slot update_slots: id 0 | task 56 | prompt processing progress, n_past = 16, n_tokens = 11, progress = 0.687500
|
||||
slot update_slots: id 0 | task 56 | prompt done, n_past = 16, n_tokens = 11
|
||||
slot release: id 0 | task 56 | stop processing: n_past = 16, truncated = 0
|
||||
slot launch_slot_: id 0 | task 59 | processing task
|
||||
slot update_slots: id 0 | task 59 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 16
|
||||
slot update_slots: id 0 | task 59 | need to evaluate at least 1 token for each active slot, n_past = 16, n_prompt_tokens = 16
|
||||
slot update_slots: id 0 | task 59 | kv cache rm [15, end)
|
||||
slot update_slots: id 0 | task 59 | prompt processing progress, n_past = 16, n_tokens = 1, progress = 0.062500
|
||||
slot update_slots: id 0 | task 59 | prompt done, n_past = 16, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 59 | stop processing: n_past = 16, truncated = 0
|
||||
slot launch_slot_: id 0 | task 57 | processing task
|
||||
slot update_slots: id 0 | task 57 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 16
|
||||
slot update_slots: id 0 | task 57 | need to evaluate at least 1 token for each active slot, n_past = 16, n_prompt_tokens = 16
|
||||
slot update_slots: id 0 | task 57 | kv cache rm [15, end)
|
||||
slot update_slots: id 0 | task 57 | prompt processing progress, n_past = 16, n_tokens = 1, progress = 0.062500
|
||||
slot update_slots: id 0 | task 57 | prompt done, n_past = 16, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 57 | stop processing: n_past = 16, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot launch_slot_: id 0 | task 62 | processing task
|
||||
slot update_slots: id 0 | task 62 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 15
|
||||
slot update_slots: id 0 | task 62 | kv cache rm [5, end)
|
||||
slot update_slots: id 0 | task 62 | prompt processing progress, n_past = 15, n_tokens = 10, progress = 0.666667
|
||||
slot update_slots: id 0 | task 62 | prompt done, n_past = 15, n_tokens = 10
|
||||
slot release: id 0 | task 62 | stop processing: n_past = 15, truncated = 0
|
||||
slot launch_slot_: id 0 | task 64 | processing task
|
||||
slot update_slots: id 0 | task 64 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 15
|
||||
slot update_slots: id 0 | task 64 | need to evaluate at least 1 token for each active slot, n_past = 15, n_prompt_tokens = 15
|
||||
slot update_slots: id 0 | task 64 | kv cache rm [14, end)
|
||||
slot update_slots: id 0 | task 64 | prompt processing progress, n_past = 15, n_tokens = 1, progress = 0.066667
|
||||
slot update_slots: id 0 | task 64 | prompt done, n_past = 15, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 64 | stop processing: n_past = 15, truncated = 0
|
||||
slot update_slots: id 0 | task 6 | new prompt, n_ctx_slot = 4096, n_keep = 0, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 6 | need to evaluate at least 1 token for each active slot, n_past = 9, n_prompt_tokens = 9
|
||||
slot update_slots: id 0 | task 6 | kv cache rm [8, end)
|
||||
slot update_slots: id 0 | task 6 | prompt processing progress, n_past = 9, n_tokens = 1, progress = 0.111111
|
||||
slot update_slots: id 0 | task 6 | prompt done, n_past = 9, n_tokens = 1
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
slot release: id 0 | task 6 | stop processing: n_past = 9, truncated = 0
|
||||
srv update_slots: all slots are idle
|
||||
srv log_server_r: request: POST /v1/embeddings 127.0.0.1 200
|
||||
|
||||
Reference in New Issue
Block a user