2025-02-02 23:01:14 +08:00
|
|
|
|
import os
|
|
|
|
|
|
import json
|
|
|
|
|
|
import uuid
|
|
|
|
|
|
import time
|
|
|
|
|
|
import queue
|
|
|
|
|
|
import asyncio
|
2025-02-24 22:04:48 +08:00
|
|
|
|
import traceback
|
2025-03-17 14:20:40 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
import threading
|
|
|
|
|
|
import websockets
|
|
|
|
|
|
from typing import Dict, Any
|
2025-03-23 16:20:20 +08:00
|
|
|
|
from plugins_func.loadplugins import auto_import_modules
|
2025-03-17 14:20:40 +08:00
|
|
|
|
from config.logger import setup_logging
|
2025-02-02 23:01:14 +08:00
|
|
|
|
from core.utils.dialogue import Message, Dialogue
|
|
|
|
|
|
from core.handle.textHandle import handleTextMessage
|
2025-03-17 14:20:40 +08:00
|
|
|
|
from core.utils.util import get_string_no_punctuation_or_emoji, extract_json_from_string, get_ip_info
|
2025-02-02 23:01:14 +08:00
|
|
|
|
from concurrent.futures import ThreadPoolExecutor, TimeoutError
|
2025-02-26 01:33:05 +08:00
|
|
|
|
from core.handle.sendAudioHandle import sendAudioMessage
|
|
|
|
|
|
from core.handle.receiveAudioHandle import handleAudioMessage
|
2025-03-15 11:48:14 +08:00
|
|
|
|
from core.handle.functionHandler import FunctionHandler
|
2025-03-20 11:52:37 +08:00
|
|
|
|
from plugins_func.register import Action, ActionResponse
|
2025-02-15 21:13:50 +08:00
|
|
|
|
from config.private_config import PrivateConfig
|
2025-02-15 19:38:32 +08:00
|
|
|
|
from core.auth import AuthMiddleware, AuthenticationError
|
2025-02-26 01:33:05 +08:00
|
|
|
|
from core.utils.auth_code_gen import AuthCodeGenerator
|
2025-03-20 08:59:45 +08:00
|
|
|
|
from core.mcp.manager import MCPManager
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-18 00:07:19 +08:00
|
|
|
|
TAG = __name__
|
|
|
|
|
|
|
2025-03-23 16:20:20 +08:00
|
|
|
|
auto_import_modules('plugins_func.functions')
|
|
|
|
|
|
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-03-04 00:35:51 +08:00
|
|
|
|
class TTSException(RuntimeError):
|
|
|
|
|
|
pass
|
|
|
|
|
|
|
|
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
class ConnectionHandler:
|
2025-03-17 02:13:10 +08:00
|
|
|
|
def __init__(self, config: Dict[str, Any], _vad, _asr, _llm, _tts, _memory, _intent):
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.config = config
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger = setup_logging()
|
2025-02-13 17:06:48 +08:00
|
|
|
|
self.auth = AuthMiddleware(config)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
self.websocket = None
|
|
|
|
|
|
self.headers = None
|
2025-03-17 14:20:40 +08:00
|
|
|
|
self.client_ip = None
|
|
|
|
|
|
self.client_ip_info = {}
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.session_id = None
|
|
|
|
|
|
self.prompt = None
|
|
|
|
|
|
self.welcome_msg = None
|
2025-02-04 13:57:05 +08:00
|
|
|
|
|
|
|
|
|
|
# 客户端状态相关
|
2025-02-04 12:20:10 +08:00
|
|
|
|
self.client_abort = False
|
2025-02-04 13:57:05 +08:00
|
|
|
|
self.client_listen_mode = "auto"
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
# 线程任务相关
|
|
|
|
|
|
self.loop = asyncio.get_event_loop()
|
|
|
|
|
|
self.stop_event = threading.Event()
|
|
|
|
|
|
self.tts_queue = queue.Queue()
|
2025-03-01 01:54:55 +08:00
|
|
|
|
self.audio_play_queue = queue.Queue()
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.executor = ThreadPoolExecutor(max_workers=10)
|
|
|
|
|
|
|
|
|
|
|
|
# 依赖的组件
|
|
|
|
|
|
self.vad = _vad
|
|
|
|
|
|
self.asr = _asr
|
|
|
|
|
|
self.llm = _llm
|
|
|
|
|
|
self.tts = _tts
|
2025-03-03 15:00:04 +08:00
|
|
|
|
self.memory = _memory
|
2025-03-09 21:33:45 +08:00
|
|
|
|
self.intent = _intent
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
# vad相关变量
|
|
|
|
|
|
self.client_audio_buffer = bytes()
|
|
|
|
|
|
self.client_have_voice = False
|
|
|
|
|
|
self.client_have_voice_last_time = 0.0
|
2025-02-16 20:42:55 +08:00
|
|
|
|
self.client_no_voice_last_time = 0.0
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.client_voice_stop = False
|
|
|
|
|
|
|
|
|
|
|
|
# asr相关变量
|
|
|
|
|
|
self.asr_audio = []
|
|
|
|
|
|
self.asr_server_receive = True
|
|
|
|
|
|
|
|
|
|
|
|
# llm相关变量
|
|
|
|
|
|
self.llm_finish_task = False
|
|
|
|
|
|
self.dialogue = Dialogue()
|
|
|
|
|
|
|
|
|
|
|
|
# tts相关变量
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.tts_first_text_index = -1
|
|
|
|
|
|
self.tts_last_text_index = -1
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-02-24 18:06:13 +08:00
|
|
|
|
# iot相关变量
|
|
|
|
|
|
self.iot_descriptors = {}
|
|
|
|
|
|
|
2025-02-14 23:09:12 +08:00
|
|
|
|
self.cmd_exit = self.config["CMD_exit"]
|
|
|
|
|
|
self.max_cmd_length = 0
|
|
|
|
|
|
for cmd in self.cmd_exit:
|
|
|
|
|
|
if len(cmd) > self.max_cmd_length:
|
|
|
|
|
|
self.max_cmd_length = len(cmd)
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-15 19:48:46 +08:00
|
|
|
|
self.private_config = None
|
2025-02-17 15:53:48 +08:00
|
|
|
|
self.auth_code_gen = AuthCodeGenerator.get_instance()
|
|
|
|
|
|
self.is_device_verified = False # 添加设备验证状态标志
|
2025-03-09 21:33:45 +08:00
|
|
|
|
self.close_after_chat = False # 是否在聊天结束后关闭连接
|
|
|
|
|
|
self.use_function_call_mode = False
|
|
|
|
|
|
if self.config["selected_module"]["Intent"] == 'function_call':
|
|
|
|
|
|
self.use_function_call_mode = True
|
2025-03-20 08:59:45 +08:00
|
|
|
|
|
|
|
|
|
|
self.mcp_manager = MCPManager(self)
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
async def handle_connection(self, ws):
|
|
|
|
|
|
try:
|
2025-02-13 17:06:48 +08:00
|
|
|
|
# 获取并验证headers
|
|
|
|
|
|
self.headers = dict(ws.request.headers)
|
2025-02-26 01:33:05 +08:00
|
|
|
|
# 获取客户端ip地址
|
2025-03-17 14:20:40 +08:00
|
|
|
|
self.client_ip = ws.remote_address[0]
|
|
|
|
|
|
self.logger.bind(tag=TAG).info(f"{self.client_ip} conn - Headers: {self.headers}")
|
2025-02-14 23:09:12 +08:00
|
|
|
|
|
2025-02-13 17:06:48 +08:00
|
|
|
|
# 进行认证
|
|
|
|
|
|
await self.auth.authenticate(self.headers)
|
2025-02-15 19:48:46 +08:00
|
|
|
|
device_id = self.headers.get("device-id", None)
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-03-23 16:20:20 +08:00
|
|
|
|
# 认证通过,继续处理
|
|
|
|
|
|
self.websocket = ws
|
|
|
|
|
|
self.session_id = str(uuid.uuid4())
|
|
|
|
|
|
|
|
|
|
|
|
self.welcome_msg = self.config["xiaozhi"]
|
|
|
|
|
|
self.welcome_msg["session_id"] = self.session_id
|
|
|
|
|
|
await self.websocket.send(json.dumps(self.welcome_msg))
|
2025-02-15 19:48:46 +08:00
|
|
|
|
# Load private configuration if device_id is provided
|
|
|
|
|
|
bUsePrivateConfig = self.config.get("use_private_config", False)
|
2025-02-18 00:45:54 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info(f"bUsePrivateConfig: {bUsePrivateConfig}, device_id: {device_id}")
|
2025-02-15 19:48:46 +08:00
|
|
|
|
if bUsePrivateConfig and device_id:
|
2025-02-17 15:53:48 +08:00
|
|
|
|
try:
|
|
|
|
|
|
self.private_config = PrivateConfig(device_id, self.config, self.auth_code_gen)
|
|
|
|
|
|
await self.private_config.load_or_create()
|
|
|
|
|
|
# 判断是否已经绑定
|
|
|
|
|
|
owner = self.private_config.get_owner()
|
|
|
|
|
|
self.is_device_verified = owner is not None
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-17 15:53:48 +08:00
|
|
|
|
if self.is_device_verified:
|
2025-02-26 01:33:05 +08:00
|
|
|
|
await self.private_config.update_last_chat_time()
|
|
|
|
|
|
|
2025-02-19 01:00:25 +08:00
|
|
|
|
llm, tts = self.private_config.create_private_instances()
|
|
|
|
|
|
if all([llm, tts]):
|
2025-02-17 15:53:48 +08:00
|
|
|
|
self.llm = llm
|
|
|
|
|
|
self.tts = tts
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info(f"Loaded private config and instances for device {device_id}")
|
2025-02-17 15:53:48 +08:00
|
|
|
|
else:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"Failed to create instances for device {device_id}")
|
2025-02-17 15:53:48 +08:00
|
|
|
|
self.private_config = None
|
|
|
|
|
|
except Exception as e:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"Error initializing private config: {e}")
|
2025-02-15 19:48:46 +08:00
|
|
|
|
self.private_config = None
|
2025-02-17 15:53:48 +08:00
|
|
|
|
raise
|
2025-02-15 19:48:46 +08:00
|
|
|
|
|
2025-03-17 14:20:40 +08:00
|
|
|
|
# 异步初始化
|
2025-03-23 16:20:20 +08:00
|
|
|
|
self.executor.submit(self._initialize_components)
|
2025-03-01 01:54:55 +08:00
|
|
|
|
# tts 消化线程
|
2025-03-24 23:53:56 +08:00
|
|
|
|
self.tts_priority_thread = threading.Thread(target=self._tts_priority_thread, daemon=True)
|
|
|
|
|
|
self.tts_priority_thread.start()
|
2025-02-13 17:06:48 +08:00
|
|
|
|
|
2025-03-01 01:54:55 +08:00
|
|
|
|
# 音频播放 消化线程
|
2025-03-24 23:53:56 +08:00
|
|
|
|
self.audio_play_priority_thread = threading.Thread(target=self._audio_play_priority_thread, daemon=True)
|
|
|
|
|
|
self.audio_play_priority_thread.start()
|
2025-03-01 01:54:55 +08:00
|
|
|
|
|
2025-02-13 17:06:48 +08:00
|
|
|
|
try:
|
|
|
|
|
|
async for message in self.websocket:
|
|
|
|
|
|
await self._route_message(message)
|
|
|
|
|
|
except websockets.exceptions.ConnectionClosed:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info("客户端断开连接")
|
2025-02-14 23:09:12 +08:00
|
|
|
|
|
2025-02-13 17:06:48 +08:00
|
|
|
|
except AuthenticationError as e:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"Authentication failed: {str(e)}")
|
2025-02-13 17:06:48 +08:00
|
|
|
|
return
|
|
|
|
|
|
except Exception as e:
|
2025-02-24 22:04:48 +08:00
|
|
|
|
stack_trace = traceback.format_exc()
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"Connection error: {str(e)}-{stack_trace}")
|
2025-02-13 17:06:48 +08:00
|
|
|
|
return
|
2025-03-03 15:00:04 +08:00
|
|
|
|
finally:
|
|
|
|
|
|
await self.memory.save_memory(self.dialogue.dialogue)
|
2025-03-24 23:53:56 +08:00
|
|
|
|
await self.close(ws)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
async def _route_message(self, message):
|
|
|
|
|
|
"""消息路由"""
|
|
|
|
|
|
if isinstance(message, str):
|
2025-02-04 13:57:05 +08:00
|
|
|
|
await handleTextMessage(self, message)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
elif isinstance(message, bytes):
|
|
|
|
|
|
await handleAudioMessage(self, message)
|
|
|
|
|
|
|
|
|
|
|
|
def _initialize_components(self):
|
2025-03-23 16:20:20 +08:00
|
|
|
|
"""加载提示词"""
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.prompt = self.config["prompt"]
|
2025-02-15 19:48:46 +08:00
|
|
|
|
if self.private_config:
|
|
|
|
|
|
self.prompt = self.private_config.private_config.get("prompt", self.prompt)
|
2025-02-05 09:06:52 +08:00
|
|
|
|
self.dialogue.put(Message(role="system", content=self.prompt))
|
2025-03-17 14:20:40 +08:00
|
|
|
|
|
2025-03-26 00:39:34 +08:00
|
|
|
|
"""加载插件"""
|
|
|
|
|
|
self.func_handler = FunctionHandler(self)
|
|
|
|
|
|
|
2025-03-23 16:20:20 +08:00
|
|
|
|
"""加载记忆"""
|
|
|
|
|
|
device_id = self.headers.get("device-id", None)
|
|
|
|
|
|
self.memory.init_memory(device_id, self.llm)
|
2025-03-25 14:27:48 +08:00
|
|
|
|
|
|
|
|
|
|
"""为意图识别设置LLM,优先使用专用LLM"""
|
|
|
|
|
|
# 检查是否配置了专用的意图识别LLM
|
2025-03-28 14:13:54 +08:00
|
|
|
|
intent_llm_name = self.config["Intent"]["intent_llm"]["llm"]
|
2025-03-25 14:27:48 +08:00
|
|
|
|
|
|
|
|
|
|
# 记录开始初始化意图识别LLM的时间
|
|
|
|
|
|
intent_llm_init_start = time.time()
|
|
|
|
|
|
|
2025-03-28 14:13:54 +08:00
|
|
|
|
if not self.use_function_call_mode and intent_llm_name and intent_llm_name in self.config["LLM"]:
|
2025-03-25 14:27:48 +08:00
|
|
|
|
# 如果配置了专用LLM,则创建独立的LLM实例
|
|
|
|
|
|
from core.utils import llm as llm_utils
|
|
|
|
|
|
intent_llm_config = self.config["LLM"][intent_llm_name]
|
|
|
|
|
|
intent_llm_type = intent_llm_config.get("type", intent_llm_name)
|
|
|
|
|
|
intent_llm = llm_utils.create_instance(intent_llm_type, intent_llm_config)
|
|
|
|
|
|
self.logger.bind(tag=TAG).info(f"为意图识别创建了专用LLM: {intent_llm_name}, 类型: {intent_llm_type}")
|
|
|
|
|
|
|
|
|
|
|
|
self.intent.set_llm(intent_llm)
|
|
|
|
|
|
else:
|
|
|
|
|
|
# 否则使用主LLM
|
|
|
|
|
|
self.intent.set_llm(self.llm)
|
|
|
|
|
|
self.logger.bind(tag=TAG).info("意图识别使用主LLM")
|
|
|
|
|
|
|
|
|
|
|
|
# 记录意图识别LLM初始化耗时
|
|
|
|
|
|
intent_llm_init_time = time.time() - intent_llm_init_start
|
|
|
|
|
|
self.logger.bind(tag=TAG).info(f"意图识别LLM初始化完成,耗时: {intent_llm_init_time:.4f}秒")
|
2025-03-23 16:20:20 +08:00
|
|
|
|
|
|
|
|
|
|
"""加载位置信息"""
|
|
|
|
|
|
self.client_ip_info = get_ip_info(self.client_ip)
|
|
|
|
|
|
if self.client_ip_info is not None and "city" in self.client_ip_info:
|
|
|
|
|
|
self.logger.bind(tag=TAG).info(f"Client ip info: {self.client_ip_info}")
|
2025-03-23 22:02:13 +08:00
|
|
|
|
self.prompt = self.prompt + f"\nuser location:{self.client_ip_info}"
|
2025-03-23 16:20:20 +08:00
|
|
|
|
self.dialogue.update_system_message(self.prompt)
|
2025-03-18 14:26:27 +08:00
|
|
|
|
|
2025-03-24 10:28:48 +08:00
|
|
|
|
"""加载MCP工具"""
|
|
|
|
|
|
asyncio.run_coroutine_threadsafe(self.mcp_manager.initialize_servers(), self.loop)
|
|
|
|
|
|
|
2025-03-15 11:48:14 +08:00
|
|
|
|
def change_system_prompt(self, prompt):
|
|
|
|
|
|
self.prompt = prompt
|
|
|
|
|
|
# 找到原来的role==system,替换原来的系统提示
|
|
|
|
|
|
for m in self.dialogue.dialogue:
|
|
|
|
|
|
if m.role == "system":
|
|
|
|
|
|
m.content = prompt
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-17 15:53:48 +08:00
|
|
|
|
async def _check_and_broadcast_auth_code(self):
|
|
|
|
|
|
"""检查设备绑定状态并广播认证码"""
|
|
|
|
|
|
if not self.private_config.get_owner():
|
|
|
|
|
|
auth_code = self.private_config.get_auth_code()
|
|
|
|
|
|
if auth_code:
|
|
|
|
|
|
# 发送验证码语音提示
|
|
|
|
|
|
text = f"请在后台输入验证码:{' '.join(auth_code)}"
|
|
|
|
|
|
self.recode_first_last_text(text)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, text)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
return False
|
|
|
|
|
|
return True
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-02-17 15:53:48 +08:00
|
|
|
|
def isNeedAuth(self):
|
|
|
|
|
|
bUsePrivateConfig = self.config.get("use_private_config", False)
|
|
|
|
|
|
if not bUsePrivateConfig:
|
|
|
|
|
|
# 如果不使用私有配置,就不需要验证
|
|
|
|
|
|
return False
|
|
|
|
|
|
return not self.is_device_verified
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
def chat(self, query):
|
2025-02-17 15:53:48 +08:00
|
|
|
|
if self.isNeedAuth():
|
|
|
|
|
|
self.llm_finish_task = True
|
2025-03-04 08:54:27 +08:00
|
|
|
|
future = asyncio.run_coroutine_threadsafe(self._check_and_broadcast_auth_code(), self.loop)
|
|
|
|
|
|
future.result()
|
2025-02-17 15:53:48 +08:00
|
|
|
|
return True
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.dialogue.put(Message(role="user", content=query))
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
response_message = []
|
2025-03-01 01:54:55 +08:00
|
|
|
|
processed_chars = 0 # 跟踪已处理的字符位置
|
2025-02-02 23:01:14 +08:00
|
|
|
|
try:
|
2025-03-01 01:54:55 +08:00
|
|
|
|
start_time = time.time()
|
2025-03-03 15:00:04 +08:00
|
|
|
|
# 使用带记忆的对话
|
2025-03-04 08:54:27 +08:00
|
|
|
|
future = asyncio.run_coroutine_threadsafe(self.memory.query_memory(query), self.loop)
|
|
|
|
|
|
memory_str = future.result()
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
|
|
|
|
|
self.logger.bind(tag=TAG).debug(f"记忆内容: {memory_str}")
|
2025-03-03 15:00:04 +08:00
|
|
|
|
llm_responses = self.llm.response(
|
2025-03-09 21:33:45 +08:00
|
|
|
|
self.session_id,
|
2025-03-03 15:00:04 +08:00
|
|
|
|
self.dialogue.get_llm_dialogue_with_memory(memory_str)
|
|
|
|
|
|
)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
except Exception as e:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"LLM 处理出错 {query}: {e}")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
return None
|
2025-03-01 01:54:55 +08:00
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.llm_finish_task = False
|
2025-03-04 00:35:51 +08:00
|
|
|
|
text_index = 0
|
2025-02-02 23:01:14 +08:00
|
|
|
|
for content in llm_responses:
|
|
|
|
|
|
response_message.append(content)
|
2025-02-04 12:20:10 +08:00
|
|
|
|
if self.client_abort:
|
|
|
|
|
|
break
|
|
|
|
|
|
|
2025-03-01 01:54:55 +08:00
|
|
|
|
end_time = time.time()
|
|
|
|
|
|
self.logger.bind(tag=TAG).debug(f"大模型返回时间: {end_time - start_time} 秒, 生成token={content}")
|
|
|
|
|
|
|
|
|
|
|
|
# 合并当前全部文本并处理未分割部分
|
|
|
|
|
|
full_text = "".join(response_message)
|
|
|
|
|
|
current_text = full_text[processed_chars:] # 从未处理的位置开始
|
|
|
|
|
|
|
|
|
|
|
|
# 查找最后一个有效标点
|
2025-03-09 21:33:45 +08:00
|
|
|
|
punctuations = ("。", "?", "!", ";", ":")
|
2025-03-01 01:54:55 +08:00
|
|
|
|
last_punct_pos = -1
|
|
|
|
|
|
for punct in punctuations:
|
|
|
|
|
|
pos = current_text.rfind(punct)
|
|
|
|
|
|
if pos > last_punct_pos:
|
|
|
|
|
|
last_punct_pos = pos
|
|
|
|
|
|
|
|
|
|
|
|
# 找到分割点则处理
|
|
|
|
|
|
if last_punct_pos != -1:
|
|
|
|
|
|
segment_text_raw = current_text[:last_punct_pos + 1]
|
|
|
|
|
|
segment_text = get_string_no_punctuation_or_emoji(segment_text_raw)
|
|
|
|
|
|
if segment_text:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
# 强制设置空字符,测试TTS出错返回语音的健壮性
|
|
|
|
|
|
# if text_index % 2 == 0:
|
|
|
|
|
|
# segment_text = " "
|
|
|
|
|
|
text_index += 1
|
|
|
|
|
|
self.recode_first_last_text(segment_text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, segment_text, text_index)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.tts_queue.put(future)
|
2025-03-01 01:54:55 +08:00
|
|
|
|
processed_chars += len(segment_text_raw) # 更新已处理字符位置
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-03-01 01:54:55 +08:00
|
|
|
|
# 处理最后剩余的文本
|
|
|
|
|
|
full_text = "".join(response_message)
|
|
|
|
|
|
remaining_text = full_text[processed_chars:]
|
|
|
|
|
|
if remaining_text:
|
|
|
|
|
|
segment_text = get_string_no_punctuation_or_emoji(remaining_text)
|
|
|
|
|
|
if segment_text:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
text_index += 1
|
|
|
|
|
|
self.recode_first_last_text(segment_text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, segment_text, text_index)
|
2025-02-14 10:29:37 +08:00
|
|
|
|
self.tts_queue.put(future)
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
self.llm_finish_task = True
|
|
|
|
|
|
self.dialogue.put(Message(role="assistant", content="".join(response_message)))
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug(json.dumps(self.dialogue.get_llm_dialogue(), indent=4, ensure_ascii=False))
|
2025-02-02 23:01:14 +08:00
|
|
|
|
return True
|
|
|
|
|
|
|
2025-03-28 17:17:14 +08:00
|
|
|
|
def chat_with_function_calling(self, query, tool_call=False):
|
2025-03-09 21:33:45 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug(f"Chat with function calling start: {query}")
|
|
|
|
|
|
"""Chat with function calling for intent detection using streaming"""
|
|
|
|
|
|
if self.isNeedAuth():
|
|
|
|
|
|
self.llm_finish_task = True
|
|
|
|
|
|
future = asyncio.run_coroutine_threadsafe(self._check_and_broadcast_auth_code(), self.loop)
|
|
|
|
|
|
future.result()
|
|
|
|
|
|
return True
|
2025-03-18 14:26:27 +08:00
|
|
|
|
|
2025-03-15 11:48:14 +08:00
|
|
|
|
if not tool_call:
|
|
|
|
|
|
self.dialogue.put(Message(role="user", content=query))
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
|
|
|
|
|
# Define intent functions
|
2025-03-25 23:48:53 +08:00
|
|
|
|
functions = None
|
|
|
|
|
|
if hasattr(self, 'func_handler'):
|
|
|
|
|
|
functions = self.func_handler.get_functions()
|
2025-03-09 21:33:45 +08:00
|
|
|
|
response_message = []
|
2025-03-28 17:17:14 +08:00
|
|
|
|
processed_chars = 0 # 跟踪已处理的字符位置
|
2025-03-18 14:26:27 +08:00
|
|
|
|
|
2025-03-09 21:33:45 +08:00
|
|
|
|
try:
|
|
|
|
|
|
start_time = time.time()
|
|
|
|
|
|
|
|
|
|
|
|
# 使用带记忆的对话
|
|
|
|
|
|
future = asyncio.run_coroutine_threadsafe(self.memory.query_memory(query), self.loop)
|
|
|
|
|
|
memory_str = future.result()
|
|
|
|
|
|
|
2025-03-28 17:17:14 +08:00
|
|
|
|
# self.logger.bind(tag=TAG).info(f"对话记录: {self.dialogue.get_llm_dialogue_with_memory(memory_str)}")
|
|
|
|
|
|
|
|
|
|
|
|
# 使用支持functions的streaming接口
|
|
|
|
|
|
llm_responses = self.llm.response_with_functions(
|
|
|
|
|
|
self.session_id,
|
|
|
|
|
|
self.dialogue.get_llm_dialogue_with_memory(memory_str),
|
|
|
|
|
|
functions=functions
|
|
|
|
|
|
)
|
2025-03-09 21:33:45 +08:00
|
|
|
|
except Exception as e:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"LLM 处理出错 {query}: {e}")
|
|
|
|
|
|
return None
|
|
|
|
|
|
|
|
|
|
|
|
self.llm_finish_task = False
|
|
|
|
|
|
text_index = 0
|
|
|
|
|
|
|
|
|
|
|
|
# 处理流式响应
|
2025-03-11 00:25:33 +08:00
|
|
|
|
tool_call_flag = False
|
|
|
|
|
|
function_name = None
|
|
|
|
|
|
function_id = None
|
|
|
|
|
|
function_arguments = ""
|
|
|
|
|
|
content_arguments = ""
|
2025-03-09 21:33:45 +08:00
|
|
|
|
for response in llm_responses:
|
2025-03-11 00:25:33 +08:00
|
|
|
|
content, tools_call = response
|
2025-03-28 17:17:14 +08:00
|
|
|
|
if "content" in response:
|
|
|
|
|
|
content = response["content"]
|
|
|
|
|
|
tools_call = None
|
2025-03-18 14:26:27 +08:00
|
|
|
|
if content is not None and len(content) > 0:
|
|
|
|
|
|
if len(response_message) <= 0 and (content == "```" or "<tool_call>" in content):
|
2025-03-11 00:25:33 +08:00
|
|
|
|
tool_call_flag = True
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
if tools_call is not None:
|
|
|
|
|
|
tool_call_flag = True
|
|
|
|
|
|
if tools_call[0].id is not None:
|
|
|
|
|
|
function_id = tools_call[0].id
|
|
|
|
|
|
if tools_call[0].function.name is not None:
|
|
|
|
|
|
function_name = tools_call[0].function.name
|
|
|
|
|
|
if tools_call[0].function.arguments is not None:
|
|
|
|
|
|
function_arguments += tools_call[0].function.arguments
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
if content is not None and len(content) > 0:
|
|
|
|
|
|
if tool_call_flag:
|
2025-03-18 14:26:27 +08:00
|
|
|
|
content_arguments += content
|
2025-03-11 00:25:33 +08:00
|
|
|
|
else:
|
|
|
|
|
|
response_message.append(content)
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
if self.client_abort:
|
|
|
|
|
|
break
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
end_time = time.time()
|
|
|
|
|
|
self.logger.bind(tag=TAG).debug(f"大模型返回时间: {end_time - start_time} 秒, 生成token={content}")
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
# 处理文本分段和TTS逻辑
|
|
|
|
|
|
# 合并当前全部文本并处理未分割部分
|
|
|
|
|
|
full_text = "".join(response_message)
|
|
|
|
|
|
current_text = full_text[processed_chars:] # 从未处理的位置开始
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
# 查找最后一个有效标点
|
|
|
|
|
|
punctuations = ("。", "?", "!", ";", ":")
|
|
|
|
|
|
last_punct_pos = -1
|
|
|
|
|
|
for punct in punctuations:
|
|
|
|
|
|
pos = current_text.rfind(punct)
|
|
|
|
|
|
if pos > last_punct_pos:
|
|
|
|
|
|
last_punct_pos = pos
|
|
|
|
|
|
|
|
|
|
|
|
# 找到分割点则处理
|
|
|
|
|
|
if last_punct_pos != -1:
|
|
|
|
|
|
segment_text_raw = current_text[:last_punct_pos + 1]
|
|
|
|
|
|
segment_text = get_string_no_punctuation_or_emoji(segment_text_raw)
|
|
|
|
|
|
if segment_text:
|
|
|
|
|
|
text_index += 1
|
|
|
|
|
|
self.recode_first_last_text(segment_text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, segment_text, text_index)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
processed_chars += len(segment_text_raw) # 更新已处理字符位置
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
2025-03-15 11:48:14 +08:00
|
|
|
|
# 处理function call
|
|
|
|
|
|
if tool_call_flag:
|
|
|
|
|
|
bHasError = False
|
|
|
|
|
|
if function_id is None:
|
|
|
|
|
|
a = extract_json_from_string(content_arguments)
|
|
|
|
|
|
if a is not None:
|
|
|
|
|
|
try:
|
|
|
|
|
|
content_arguments_json = json.loads(a)
|
|
|
|
|
|
function_name = content_arguments_json["name"]
|
|
|
|
|
|
function_arguments = json.dumps(content_arguments_json["arguments"], ensure_ascii=False)
|
|
|
|
|
|
function_id = str(uuid.uuid4().hex)
|
|
|
|
|
|
except Exception as e:
|
|
|
|
|
|
bHasError = True
|
|
|
|
|
|
response_message.append(a)
|
|
|
|
|
|
else:
|
|
|
|
|
|
bHasError = True
|
|
|
|
|
|
response_message.append(content_arguments)
|
|
|
|
|
|
if bHasError:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"function call error: {content_arguments}")
|
|
|
|
|
|
else:
|
|
|
|
|
|
function_arguments = json.loads(function_arguments)
|
|
|
|
|
|
if not bHasError:
|
2025-03-18 14:26:27 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info(
|
|
|
|
|
|
f"function_name={function_name}, function_id={function_id}, function_arguments={function_arguments}")
|
2025-03-15 11:48:14 +08:00
|
|
|
|
function_call_data = {
|
|
|
|
|
|
"name": function_name,
|
|
|
|
|
|
"id": function_id,
|
|
|
|
|
|
"arguments": function_arguments
|
|
|
|
|
|
}
|
2025-03-20 08:59:45 +08:00
|
|
|
|
|
|
|
|
|
|
# 处理MCP工具调用
|
|
|
|
|
|
if self.mcp_manager.is_mcp_tool(function_name):
|
2025-03-20 11:52:37 +08:00
|
|
|
|
result = self._handle_mcp_tool_call(function_call_data)
|
2025-03-20 08:59:45 +08:00
|
|
|
|
else:
|
|
|
|
|
|
# 处理系统函数
|
|
|
|
|
|
result = self.func_handler.handle_llm_function_call(self, function_call_data)
|
2025-03-20 11:52:37 +08:00
|
|
|
|
self._handle_function_result(result, function_call_data, text_index+1)
|
2025-03-15 11:48:14 +08:00
|
|
|
|
|
2025-03-18 14:26:27 +08:00
|
|
|
|
# 处理最后剩余的文本
|
2025-03-09 21:33:45 +08:00
|
|
|
|
full_text = "".join(response_message)
|
|
|
|
|
|
remaining_text = full_text[processed_chars:]
|
|
|
|
|
|
if remaining_text:
|
|
|
|
|
|
segment_text = get_string_no_punctuation_or_emoji(remaining_text)
|
|
|
|
|
|
if segment_text:
|
|
|
|
|
|
text_index += 1
|
|
|
|
|
|
self.recode_first_last_text(segment_text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, segment_text, text_index)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
|
|
|
|
|
|
# 存储对话内容
|
2025-03-18 14:26:27 +08:00
|
|
|
|
if len(response_message) > 0:
|
2025-03-11 00:25:33 +08:00
|
|
|
|
self.dialogue.put(Message(role="assistant", content="".join(response_message)))
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
|
|
|
|
|
self.llm_finish_task = True
|
|
|
|
|
|
self.logger.bind(tag=TAG).debug(json.dumps(self.dialogue.get_llm_dialogue(), indent=4, ensure_ascii=False))
|
|
|
|
|
|
|
|
|
|
|
|
return True
|
|
|
|
|
|
|
2025-03-20 11:52:37 +08:00
|
|
|
|
def _handle_mcp_tool_call(self, function_call_data):
|
|
|
|
|
|
function_arguments = function_call_data["arguments"]
|
|
|
|
|
|
function_name = function_call_data["name"]
|
|
|
|
|
|
try:
|
|
|
|
|
|
args_dict = function_arguments
|
|
|
|
|
|
if isinstance(function_arguments, str):
|
|
|
|
|
|
try:
|
|
|
|
|
|
args_dict = json.loads(function_arguments)
|
|
|
|
|
|
except json.JSONDecodeError:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"无法解析 function_arguments: {function_arguments}")
|
|
|
|
|
|
return ActionResponse(action=Action.REQLLM, result="参数解析失败", response="")
|
|
|
|
|
|
|
|
|
|
|
|
tool_result = asyncio.run_coroutine_threadsafe(self.mcp_manager.execute_tool(
|
|
|
|
|
|
function_name,
|
|
|
|
|
|
args_dict
|
|
|
|
|
|
), self.loop).result()
|
|
|
|
|
|
# meta=None content=[TextContent(type='text', text='北京当前天气:\n温度: 21°C\n天气: 晴\n湿度: 6%\n风向: 西北 风\n风力等级: 5级', annotations=None)] isError=False
|
2025-03-20 18:20:09 +08:00
|
|
|
|
content_text = ""
|
|
|
|
|
|
if tool_result is not None and tool_result.content is not None:
|
|
|
|
|
|
for content in tool_result.content:
|
|
|
|
|
|
content_type = content.type
|
|
|
|
|
|
if content_type == "text":
|
|
|
|
|
|
content_text = content.text
|
|
|
|
|
|
elif content_type == "image":
|
|
|
|
|
|
pass
|
|
|
|
|
|
|
|
|
|
|
|
if len(content_text) > 0:
|
|
|
|
|
|
return ActionResponse(action=Action.REQLLM, result=content_text, response="")
|
2025-03-20 11:52:37 +08:00
|
|
|
|
|
|
|
|
|
|
except Exception as e:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"MCP工具调用错误: {e}")
|
|
|
|
|
|
return ActionResponse(action=Action.REQLLM, result="工具调用出错", response="")
|
|
|
|
|
|
|
|
|
|
|
|
return ActionResponse(action=Action.REQLLM, result="工具调用出错", response="")
|
|
|
|
|
|
|
|
|
|
|
|
|
2025-03-11 00:25:33 +08:00
|
|
|
|
def _handle_function_result(self, result, function_call_data, text_index):
|
2025-03-18 14:26:27 +08:00
|
|
|
|
if result.action == Action.RESPONSE: # 直接回复前端
|
2025-03-11 00:25:33 +08:00
|
|
|
|
text = result.response
|
|
|
|
|
|
self.recode_first_last_text(text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, text, text_index)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
self.dialogue.put(Message(role="assistant", content=text))
|
2025-03-18 14:26:27 +08:00
|
|
|
|
elif result.action == Action.REQLLM: # 调用函数后再请求llm生成回复
|
2025-03-15 11:48:14 +08:00
|
|
|
|
|
2025-03-28 17:17:14 +08:00
|
|
|
|
text = result.result
|
|
|
|
|
|
if text is not None and len(text) > 0:
|
|
|
|
|
|
function_id = function_call_data["id"]
|
|
|
|
|
|
function_name = function_call_data["name"]
|
|
|
|
|
|
function_arguments = function_call_data["arguments"]
|
|
|
|
|
|
self.dialogue.put(Message(role='assistant',
|
|
|
|
|
|
tool_calls=[{"id": function_id,
|
|
|
|
|
|
"function": {"arguments": function_arguments,
|
|
|
|
|
|
"name": function_name},
|
|
|
|
|
|
"type": 'function',
|
|
|
|
|
|
"index": 0}]))
|
|
|
|
|
|
|
|
|
|
|
|
self.dialogue.put(Message(role="tool", tool_call_id=function_id, content=text))
|
|
|
|
|
|
self.chat_with_function_calling(text, tool_call=True)
|
2025-03-18 00:34:16 +08:00
|
|
|
|
elif result.action == Action.NOTFOUND:
|
2025-03-18 14:26:27 +08:00
|
|
|
|
text = result.result
|
|
|
|
|
|
self.recode_first_last_text(text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, text, text_index)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
self.dialogue.put(Message(role="assistant", content=text))
|
2025-03-18 00:34:16 +08:00
|
|
|
|
else:
|
2025-03-18 14:26:27 +08:00
|
|
|
|
text = result.result
|
|
|
|
|
|
self.recode_first_last_text(text, text_index)
|
|
|
|
|
|
future = self.executor.submit(self.speak_and_play, text, text_index)
|
|
|
|
|
|
self.tts_queue.put(future)
|
|
|
|
|
|
self.dialogue.put(Message(role="assistant", content=text))
|
2025-03-11 00:25:33 +08:00
|
|
|
|
|
2025-03-01 01:54:55 +08:00
|
|
|
|
def _tts_priority_thread(self):
|
2025-02-02 23:01:14 +08:00
|
|
|
|
while not self.stop_event.is_set():
|
|
|
|
|
|
text = None
|
|
|
|
|
|
try:
|
2025-03-24 23:53:56 +08:00
|
|
|
|
try:
|
|
|
|
|
|
future = self.tts_queue.get(timeout=1)
|
|
|
|
|
|
except queue.Empty:
|
|
|
|
|
|
if self.stop_event.is_set():
|
|
|
|
|
|
break
|
|
|
|
|
|
continue
|
2025-02-11 12:48:39 +08:00
|
|
|
|
if future is None:
|
|
|
|
|
|
continue
|
2025-02-02 23:01:14 +08:00
|
|
|
|
text = None
|
2025-03-04 00:35:51 +08:00
|
|
|
|
opus_datas, text_index, tts_file = [], 0, None
|
2025-02-02 23:01:14 +08:00
|
|
|
|
try:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug("正在处理TTS任务...")
|
2025-03-15 00:19:45 +08:00
|
|
|
|
tts_timeout = self.config.get("tts_timeout", 10)
|
|
|
|
|
|
tts_file, text, text_index = future.result(timeout=tts_timeout)
|
2025-02-14 10:29:37 +08:00
|
|
|
|
if text is None or len(text) <= 0:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"TTS出错:{text_index}: tts text is empty")
|
|
|
|
|
|
elif tts_file is None:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"TTS出错: file is empty: {text_index}: {text}")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
else:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug(f"TTS生成:文件路径: {tts_file}")
|
|
|
|
|
|
if os.path.exists(tts_file):
|
2025-03-15 00:19:45 +08:00
|
|
|
|
opus_datas, duration = self.tts.audio_to_opus_data(tts_file)
|
2025-03-04 00:35:51 +08:00
|
|
|
|
else:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"TTS出错:文件不存在{tts_file}")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
except TimeoutError:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error("TTS超时")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
except Exception as e:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"TTS出错: {e}")
|
2025-02-04 12:20:10 +08:00
|
|
|
|
if not self.client_abort:
|
|
|
|
|
|
# 如果没有中途打断就发送语音
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.audio_play_queue.put((opus_datas, text, text_index))
|
|
|
|
|
|
if self.tts.delete_audio_file and tts_file is not None and os.path.exists(tts_file):
|
2025-02-02 23:01:14 +08:00
|
|
|
|
os.remove(tts_file)
|
|
|
|
|
|
except Exception as e:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"TTS任务处理错误: {e}")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.clearSpeakStatus()
|
|
|
|
|
|
asyncio.run_coroutine_threadsafe(
|
|
|
|
|
|
self.websocket.send(json.dumps({"type": "tts", "state": "stop", "session_id": self.session_id})),
|
|
|
|
|
|
self.loop
|
|
|
|
|
|
)
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"tts_priority priority_thread: {text} {e}")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-03-01 01:54:55 +08:00
|
|
|
|
def _audio_play_priority_thread(self):
|
|
|
|
|
|
while not self.stop_event.is_set():
|
|
|
|
|
|
text = None
|
|
|
|
|
|
try:
|
2025-03-24 23:53:56 +08:00
|
|
|
|
try:
|
|
|
|
|
|
opus_datas, text, text_index = self.audio_play_queue.get(timeout=1)
|
|
|
|
|
|
except queue.Empty:
|
|
|
|
|
|
if self.stop_event.is_set():
|
|
|
|
|
|
break
|
|
|
|
|
|
continue
|
2025-03-04 00:35:51 +08:00
|
|
|
|
future = asyncio.run_coroutine_threadsafe(sendAudioMessage(self, opus_datas, text, text_index),
|
|
|
|
|
|
self.loop)
|
2025-03-01 01:54:55 +08:00
|
|
|
|
future.result()
|
|
|
|
|
|
except Exception as e:
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"audio_play_priority priority_thread: {text} {e}")
|
2025-03-01 01:54:55 +08:00
|
|
|
|
|
2025-03-04 00:35:51 +08:00
|
|
|
|
def speak_and_play(self, text, text_index=0):
|
2025-02-02 23:01:14 +08:00
|
|
|
|
if text is None or len(text) <= 0:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info(f"无需tts转换,query为空,{text}")
|
2025-03-04 00:35:51 +08:00
|
|
|
|
return None, text, text_index
|
2025-02-02 23:01:14 +08:00
|
|
|
|
tts_file = self.tts.to_tts(text)
|
|
|
|
|
|
if tts_file is None:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).error(f"tts转换失败,{text}")
|
2025-03-04 00:35:51 +08:00
|
|
|
|
return None, text, text_index
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug(f"TTS 文件生成完毕: {tts_file}")
|
2025-03-04 00:35:51 +08:00
|
|
|
|
return tts_file, text, text_index
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
|
|
|
|
|
def clearSpeakStatus(self):
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug(f"清除服务端讲话状态")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
self.asr_server_receive = True
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.tts_last_text_index = -1
|
|
|
|
|
|
self.tts_first_text_index = -1
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-03-04 00:35:51 +08:00
|
|
|
|
def recode_first_last_text(self, text, text_index=0):
|
|
|
|
|
|
if self.tts_first_text_index == -1:
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info(f"大模型说出第一句话: {text}")
|
2025-03-04 00:35:51 +08:00
|
|
|
|
self.tts_first_text_index = text_index
|
|
|
|
|
|
self.tts_last_text_index = text_index
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-03-24 23:53:56 +08:00
|
|
|
|
async def close(self, ws=None):
|
2025-02-02 23:01:14 +08:00
|
|
|
|
"""资源清理方法"""
|
2025-03-20 08:59:45 +08:00
|
|
|
|
# 清理MCP资源
|
|
|
|
|
|
await self.mcp_manager.cleanup_all()
|
2025-02-26 01:33:05 +08:00
|
|
|
|
|
2025-03-24 23:53:56 +08:00
|
|
|
|
# 触发停止事件并清理资源
|
|
|
|
|
|
if self.stop_event:
|
|
|
|
|
|
self.stop_event.set()
|
|
|
|
|
|
|
|
|
|
|
|
# 立即关闭线程池
|
|
|
|
|
|
if self.executor:
|
|
|
|
|
|
self.executor.shutdown(wait=False, cancel_futures=True)
|
|
|
|
|
|
self.executor = None
|
|
|
|
|
|
|
|
|
|
|
|
# 清空任务队列
|
|
|
|
|
|
self._clear_queues()
|
|
|
|
|
|
|
|
|
|
|
|
if ws:
|
|
|
|
|
|
await ws.close()
|
|
|
|
|
|
elif self.websocket:
|
2025-02-02 23:01:14 +08:00
|
|
|
|
await self.websocket.close()
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).info("连接资源已释放")
|
2025-02-02 23:01:14 +08:00
|
|
|
|
|
2025-03-24 23:53:56 +08:00
|
|
|
|
def _clear_queues(self):
|
|
|
|
|
|
# 清空所有任务队列
|
|
|
|
|
|
for q in [self.tts_queue, self.audio_play_queue]:
|
|
|
|
|
|
if not q:
|
|
|
|
|
|
continue
|
|
|
|
|
|
while not q.empty():
|
|
|
|
|
|
try:
|
|
|
|
|
|
q.get_nowait()
|
|
|
|
|
|
except queue.Empty:
|
|
|
|
|
|
continue
|
|
|
|
|
|
q.queue.clear()
|
|
|
|
|
|
# 添加毒丸信号到队列,确保线程退出
|
|
|
|
|
|
# q.queue.put(None)
|
|
|
|
|
|
|
2025-02-02 23:01:14 +08:00
|
|
|
|
def reset_vad_states(self):
|
|
|
|
|
|
self.client_audio_buffer = bytes()
|
|
|
|
|
|
self.client_have_voice = False
|
|
|
|
|
|
self.client_have_voice_last_time = 0
|
|
|
|
|
|
self.client_voice_stop = False
|
2025-02-18 00:07:19 +08:00
|
|
|
|
self.logger.bind(tag=TAG).debug("VAD states reset.")
|
2025-03-09 21:33:45 +08:00
|
|
|
|
|
|
|
|
|
|
def chat_and_close(self, text):
|
|
|
|
|
|
"""Chat with the user and then close the connection"""
|
|
|
|
|
|
try:
|
|
|
|
|
|
# Use the existing chat method
|
|
|
|
|
|
self.chat(text)
|
|
|
|
|
|
|
|
|
|
|
|
# After chat is complete, close the connection
|
|
|
|
|
|
self.close_after_chat = True
|
|
|
|
|
|
except Exception as e:
|
|
|
|
|
|
self.logger.bind(tag=TAG).error(f"Chat and close error: {str(e)}")
|