Source code for langchain_community.utilities.google_trends

"""调用谷歌学术搜索的工具。"""
from typing import Any, Dict, Optional, cast

from langchain_core.pydantic_v1 import BaseModel, Extra, SecretStr, root_validator
from langchain_core.utils import convert_to_secret_str, get_from_dict_or_env


[docs]class GoogleTrendsAPIWrapper(BaseModel): """封装了SerpApi的Google Scholar API 您可以通过在以下网址注册来创建SerpApi.com密钥:https://serpapi.com/users/sign_up。 该封装使用了SerpApi.com的Python包: https://serpapi.com/integrations/python 要使用,您应该设置环境变量``SERPAPI_API_KEY`` 为您的API密钥,或将`serp_api_key`作为命名参数传递给构造函数。 示例: .. code-block:: python from langchain_community.utilities import GoogleTrendsAPIWrapper google_trends = GoogleTrendsAPIWrapper() google_trends.run('langchain') """ serp_search_engine: Any serp_api_key: Optional[SecretStr] = None class Config: """此pydantic对象的配置。""" extra = Extra.forbid @root_validator() def validate_environment(cls, values: Dict) -> Dict: """验证环境中是否存在API密钥和Python包。""" values["serp_api_key"] = convert_to_secret_str( get_from_dict_or_env(values, "serp_api_key", "SERPAPI_API_KEY") ) try: from serpapi import SerpApiClient except ImportError: raise ImportError( "google-search-results is not installed. " "Please install it with `pip install google-search-results" ">=2.4.2`" ) serp_search_engine = SerpApiClient values["serp_search_engine"] = serp_search_engine return values
[docs] def run(self, query: str) -> str: """通过Serpapi在Google Trends中运行查询。""" serpapi_api_key = cast(SecretStr, self.serp_api_key) params = { "engine": "google_trends", "api_key": serpapi_api_key.get_secret_value(), "q": query, } total_results = [] client = self.serp_search_engine(params) client_dict = client.get_dict() total_results = ( client_dict["interest_over_time"]["timeline_data"] if "interest_over_time" in client_dict else None ) if not total_results: return "No good Trend Result was found" start_date = total_results[0]["date"].split() end_date = total_results[-1]["date"].split() values = [ results.get("values")[0].get("extracted_value") for results in total_results ] min_value = min(values) max_value = max(values) avg_value = sum(values) / len(values) percentage_change = ( (values[-1] - values[0]) / (values[0] if values[0] != 0 else 1) * (100 if values[0] != 0 else 1) ) params = { "engine": "google_trends", "api_key": serpapi_api_key.get_secret_value(), "data_type": "RELATED_QUERIES", "q": query, } total_results2 = {} client = self.serp_search_engine(params) total_results2 = client.get_dict().get("related_queries", {}) rising = [] top = [] rising = [results.get("query") for results in total_results2.get("rising", [])] top = [results.get("query") for results in total_results2.get("top", [])] doc = [ f"Query: {query}\n" f"Date From: {start_date[0]} {start_date[1]}, {start_date[-1]}\n" f"Date To: {end_date[0]} {end_date[3]} {end_date[-1]}\n" f"Min Value: {min_value}\n" f"Max Value: {max_value}\n" f"Average Value: {avg_value}\n" f"Percent Change: {str(percentage_change) + '%'}\n" f"Trend values: {', '.join([str(x) for x in values])}\n" f"Rising Related Queries: {', '.join(rising)}\n" f"Top Related Queries: {', '.join(top)}" ] return "\n\n".join(doc)