Vous ne pouvez pas sélectionner plus de 25 sujets Les noms de sujets doivent commencer par une lettre ou un nombre, peuvent contenir des tirets ('-') et peuvent comporter jusqu'à 35 caractères.

file_service.py 8.4KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243
  1. #
  2. # Copyright 2024 The InfiniFlow Authors. All Rights Reserved.
  3. #
  4. # Licensed under the Apache License, Version 2.0 (the "License");
  5. # you may not use this file except in compliance with the License.
  6. # You may obtain a copy of the License at
  7. #
  8. # http://www.apache.org/licenses/LICENSE-2.0
  9. #
  10. # Unless required by applicable law or agreed to in writing, software
  11. # distributed under the License is distributed on an "AS IS" BASIS,
  12. # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied.
  13. # See the License for the specific language governing permissions and
  14. # limitations under the License.
  15. #
  16. from flask_login import current_user
  17. from peewee import fn
  18. from api.db import FileType
  19. from api.db.db_models import DB, File2Document, Knowledgebase
  20. from api.db.db_models import File, Document
  21. from api.db.services.common_service import CommonService
  22. from api.utils import get_uuid
  23. class FileService(CommonService):
  24. model = File
  25. @classmethod
  26. @DB.connection_context()
  27. def get_by_pf_id(cls, tenant_id, pf_id, page_number, items_per_page,
  28. orderby, desc, keywords):
  29. if keywords:
  30. files = cls.model.select().where(
  31. (cls.model.tenant_id == tenant_id)
  32. & (cls.model.parent_id == pf_id), (fn.LOWER(cls.model.name).like(f"%%{keywords.lower()}%%")))
  33. else:
  34. files = cls.model.select().where((cls.model.tenant_id == tenant_id)
  35. & (cls.model.parent_id == pf_id))
  36. count = files.count()
  37. if desc:
  38. files = files.order_by(cls.model.getter_by(orderby).desc())
  39. else:
  40. files = files.order_by(cls.model.getter_by(orderby).asc())
  41. files = files.paginate(page_number, items_per_page)
  42. res_files = list(files.dicts())
  43. for file in res_files:
  44. if file["type"] == FileType.FOLDER.value:
  45. file["size"] = cls.get_folder_size(file["id"])
  46. file['kbs_info'] = []
  47. continue
  48. kbs_info = cls.get_kb_id_by_file_id(file['id'])
  49. file['kbs_info'] = kbs_info
  50. return res_files, count
  51. @classmethod
  52. @DB.connection_context()
  53. def get_kb_id_by_file_id(cls, file_id):
  54. kbs = (cls.model.select(*[Knowledgebase.id, Knowledgebase.name])
  55. .join(File2Document, on=(File2Document.file_id == file_id))
  56. .join(Document, on=(File2Document.document_id == Document.id))
  57. .join(Knowledgebase, on=(Knowledgebase.id == Document.kb_id))
  58. .where(cls.model.id == file_id))
  59. if not kbs: return []
  60. kbs_info_list = []
  61. for kb in list(kbs.dicts()):
  62. kbs_info_list.append({"kb_id": kb['id'], "kb_name": kb['name']})
  63. return kbs_info_list
  64. @classmethod
  65. @DB.connection_context()
  66. def get_by_pf_id_name(cls, id, name):
  67. file = cls.model.select().where((cls.model.parent_id == id) & (cls.model.name == name))
  68. if file.count():
  69. e, file = cls.get_by_id(file[0].id)
  70. if not e:
  71. raise RuntimeError("Database error (File retrieval)!")
  72. return file
  73. return None
  74. @classmethod
  75. @DB.connection_context()
  76. def get_id_list_by_id(cls, id, name, count, res):
  77. if count < len(name):
  78. file = cls.get_by_pf_id_name(id, name[count])
  79. if file:
  80. res.append(file.id)
  81. return cls.get_id_list_by_id(file.id, name, count + 1, res)
  82. else:
  83. return res
  84. else:
  85. return res
  86. @classmethod
  87. @DB.connection_context()
  88. def get_all_innermost_file_ids(cls, folder_id, result_ids):
  89. subfolders = cls.model.select().where(cls.model.parent_id == folder_id)
  90. if subfolders.exists():
  91. for subfolder in subfolders:
  92. cls.get_all_innermost_file_ids(subfolder.id, result_ids)
  93. else:
  94. result_ids.append(folder_id)
  95. return result_ids
  96. @classmethod
  97. @DB.connection_context()
  98. def create_folder(cls, file, parent_id, name, count):
  99. if count > len(name) - 2:
  100. return file
  101. else:
  102. file = cls.insert({
  103. "id": get_uuid(),
  104. "parent_id": parent_id,
  105. "tenant_id": current_user.id,
  106. "created_by": current_user.id,
  107. "name": name[count],
  108. "location": "",
  109. "size": 0,
  110. "type": FileType.FOLDER.value
  111. })
  112. return cls.create_folder(file, file.id, name, count + 1)
  113. @classmethod
  114. @DB.connection_context()
  115. def is_parent_folder_exist(cls, parent_id):
  116. parent_files = cls.model.select().where(cls.model.id == parent_id)
  117. if parent_files.count():
  118. return True
  119. cls.delete_folder_by_pf_id(parent_id)
  120. return False
  121. @classmethod
  122. @DB.connection_context()
  123. def get_root_folder(cls, tenant_id):
  124. file = cls.model.select().where(cls.model.tenant_id == tenant_id and
  125. cls.model.parent_id == cls.model.id)
  126. if not file:
  127. file_id = get_uuid()
  128. file = {
  129. "id": file_id,
  130. "parent_id": file_id,
  131. "tenant_id": tenant_id,
  132. "created_by": tenant_id,
  133. "name": "/",
  134. "type": FileType.FOLDER.value,
  135. "size": 0,
  136. "location": "",
  137. }
  138. cls.save(**file)
  139. else:
  140. file_id = file[0].id
  141. e, file = cls.get_by_id(file_id)
  142. if not e:
  143. raise RuntimeError("Database error (File retrieval)!")
  144. return file
  145. @classmethod
  146. @DB.connection_context()
  147. def get_parent_folder(cls, file_id):
  148. file = cls.model.select().where(cls.model.id == file_id)
  149. if file.count():
  150. e, file = cls.get_by_id(file[0].parent_id)
  151. if not e:
  152. raise RuntimeError("Database error (File retrieval)!")
  153. else:
  154. raise RuntimeError("Database error (File doesn't exist)!")
  155. return file
  156. @classmethod
  157. @DB.connection_context()
  158. def get_all_parent_folders(cls, start_id):
  159. parent_folders = []
  160. current_id = start_id
  161. while current_id:
  162. e, file = cls.get_by_id(current_id)
  163. if file.parent_id != file.id and e:
  164. parent_folders.append(file)
  165. current_id = file.parent_id
  166. else:
  167. parent_folders.append(file)
  168. break
  169. return parent_folders
  170. @classmethod
  171. @DB.connection_context()
  172. def insert(cls, file):
  173. if not cls.save(**file):
  174. raise RuntimeError("Database error (File)!")
  175. e, file = cls.get_by_id(file["id"])
  176. if not e:
  177. raise RuntimeError("Database error (File retrieval)!")
  178. return file
  179. @classmethod
  180. @DB.connection_context()
  181. def delete(cls, file):
  182. return cls.delete_by_id(file.id)
  183. @classmethod
  184. @DB.connection_context()
  185. def delete_by_pf_id(cls, folder_id):
  186. return cls.model.delete().where(cls.model.parent_id == folder_id).execute()
  187. @classmethod
  188. @DB.connection_context()
  189. def delete_folder_by_pf_id(cls, user_id, folder_id):
  190. try:
  191. files = cls.model.select().where((cls.model.tenant_id == user_id)
  192. & (cls.model.parent_id == folder_id))
  193. for file in files:
  194. cls.delete_folder_by_pf_id(user_id, file.id)
  195. return cls.model.delete().where((cls.model.tenant_id == user_id)
  196. & (cls.model.id == folder_id)).execute(),
  197. except Exception as e:
  198. print(e)
  199. raise RuntimeError("Database error (File retrieval)!")
  200. @classmethod
  201. @DB.connection_context()
  202. def get_file_count(cls, tenant_id):
  203. files = cls.model.select(cls.model.id).where(cls.model.tenant_id == tenant_id)
  204. return len(files)
  205. @classmethod
  206. @DB.connection_context()
  207. def get_folder_size(cls, folder_id):
  208. size = 0
  209. def dfs(parent_id):
  210. nonlocal size
  211. for f in cls.model.select(*[cls.model.id, cls.model.size, cls.model.type]).where(
  212. cls.model.parent_id == parent_id, cls.model.id != parent_id):
  213. size += f.size
  214. if f.type == FileType.FOLDER.value:
  215. dfs(f.id)
  216. dfs(folder_id)
  217. return size