Refa: PARALLEL_DEVICES is a static parameter. (#6168)

### What problem does this PR solve? ### Type of change - [x] Refactoring
2026-01-23 03:26:53 +08:00 · 2025-03-17 16:49:54 +08:00
parent 45fe02c8b3
commit 3a99c2b5f4
6 changed files with 29 additions and 28 deletions
--- a/rag/svr/task_executor.py
+++ b/rag/svr/task_executor.py
@ -100,13 +100,6 @@ MAX_CONCURRENT_CHUNK_BUILDERS = int(os.environ.get('MAX_CONCURRENT_CHUNK_BUILDER
 task_limiter = trio.CapacityLimiter(MAX_CONCURRENT_TASKS)
 chunk_limiter = trio.CapacityLimiter(MAX_CONCURRENT_CHUNK_BUILDERS)

-PARALLEL_DEVICES = None
-try:
-    import torch.cuda
-    PARALLEL_DEVICES = torch.cuda.device_count()
-    logging.info(f"found {PARALLEL_DEVICES} gpus")
-except Exception:
-    logging.info("can't import package 'torch'")

 # SIGUSR1 handler: start tracemalloc and take snapshot
 def start_tracemalloc_and_snapshot(signum, frame):
@ -249,7 +242,7 @@ async def build_chunks(task, progress_callback):
    try:
        async with chunk_limiter:
            cks = await trio.to_thread.run_sync(lambda: chunker.chunk(task["name"], binary=binary, from_page=task["from_page"],
-                                to_page=task["to_page"], lang=task["language"], parallel_devices = PARALLEL_DEVICES, callback=progress_callback,
+                                to_page=task["to_page"], lang=task["language"], callback=progress_callback,
                                kb_id=task["kb_id"], parser_config=task["parser_config"], tenant_id=task["tenant_id"]))
        logging.info("Chunking({}) {}/{} done".format(timer() - st, task["location"], task["name"]))
    except TaskCanceledException: