
Signed-off-by: yihong0618 <zouzou0208@gmail.com> Signed-off-by: -LAN- <laipz8200@outlook.com> Co-authored-by: AkaraChen <akarachen@outlook.com> Co-authored-by: nite-knite <nkCoding@gmail.com> Co-authored-by: Joel <iamjoel007@gmail.com> Co-authored-by: Warren Chen <warren.chen830@gmail.com> Co-authored-by: crazywoola <427733928@qq.com> Co-authored-by: Yi Xiao <54782454+YIXIAO0@users.noreply.github.com> Co-authored-by: yihong <zouzou0208@gmail.com> Co-authored-by: -LAN- <laipz8200@outlook.com> Co-authored-by: KVOJJJin <jzongcode@gmail.com> Co-authored-by: github-actions[bot] <41898282+github-actions[bot]@users.noreply.github.com> Co-authored-by: JzoNgKVO <27049666+JzoNgKVO@users.noreply.github.com> Co-authored-by: Charlie.Wei <luowei@cvte.com> Co-authored-by: crazywoola <100913391+crazywoola@users.noreply.github.com> Co-authored-by: huayaoyue6 <huayaoyue@163.com> Co-authored-by: kurokobo <kuro664@gmail.com> Co-authored-by: Matsuda <yiyth.fcb6@gmail.com> Co-authored-by: shirochan <s.yusuke0711@gmail.com> Co-authored-by: Jyong <76649700+JohnJyong@users.noreply.github.com> Co-authored-by: Huỳnh Gia Bôi <boihuynh147@gmail.com> Co-authored-by: Julian Huynh <julian.huynh@immersio.io> Co-authored-by: Hash Brown <hi@xzd.me> Co-authored-by: 非法操作 <hjlarry@163.com> Co-authored-by: Kazuki Takamatsu <kazuki.takamatsu@chowagiken.co.jp> Co-authored-by: Trey Dong <1346650911@qq.com> Co-authored-by: VoidIsVoid <343750470@qq.com> Co-authored-by: Gimling <huangjl@ruyi.ai> Co-authored-by: xiandan-erizo <xiandan.erizo@gmail.com> Co-authored-by: Muneyuki Noguchi <nogu.dev@gmail.com> Co-authored-by: zhaobingshuang <1475195565@qq.com> Co-authored-by: zhaobs <zhaobs@cailian.net> Co-authored-by: suzuki.sh <s2terminal@users.noreply.github.com> Co-authored-by: Yingchun Lai <laiyingchun@apache.org> Co-authored-by: huanshare <huanshare@live.com> Co-authored-by: huanshare <liuhuan101@longfor.com> Co-authored-by: orangeclk <orangeclk@users.noreply.github.com> Co-authored-by: 문정현 <120004247+JungHyunMoon@users.noreply.github.com> Co-authored-by: barabicu <kztk533@gmail.com> Co-authored-by: Wei Mingzhi <whistler_wmz@users.sf.net> Co-authored-by: Paul van Oorschot <20116814+pvoo@users.noreply.github.com> Co-authored-by: zkyTech <zhangkunyuan@hotmail.com> Co-authored-by: zhangkunyuan <zhangkunyuan@cmhi.chinamobile.com> Co-authored-by: Tommy <34446820+Asterovim@users.noreply.github.com> Co-authored-by: zxhlyh <jasonapring2015@outlook.com> Co-authored-by: Novice <857526207@qq.com> Co-authored-by: Novice Lee <novicelee@NovicedeMacBook-Pro.local> Co-authored-by: Novice Lee <novicelee@NoviPro.local> Co-authored-by: zxhlyh <16177003+zxhlyh@users.noreply.github.com> Co-authored-by: liuzhenghua <1090179900@qq.com> Co-authored-by: Jiang <65766008+AlwaysBluer@users.noreply.github.com> Co-authored-by: jiangzhijie <jiangzhijie.jzj@alibaba-inc.com> Co-authored-by: Joe <79627742+ZhouhaoJiang@users.noreply.github.com> Co-authored-by: Alok Shrivastwa <alok.shrivastwa@gmail.com> Co-authored-by: Alok Shrivastwa <Alok.Shrivastwa@microland.com> Co-authored-by: JasonVV <jasonwangiii@outlook.com> Co-authored-by: Hiroshi Fujita <fujita-h@users.noreply.github.com> Co-authored-by: Kevin9703 <51311316+Kevin9703@users.noreply.github.com> Co-authored-by: NFish <douxc512@gmail.com> Co-authored-by: Junyan Qin <1010553892@qq.com> Co-authored-by: IWAI, Masaharu <iwaim.sub@gmail.com> Co-authored-by: IWAI, Masaharu <iwai_masaharu@funkit.co.jp> Co-authored-by: Bowen Liang <liangbowen@gf.com.cn> Co-authored-by: luckylhb90 <luckylhb90@gmail.com> Co-authored-by: hobo.l <hobo.l@binance.com> Co-authored-by: douxc <7553076+douxc@users.noreply.github.com>
178 lines
6.4 KiB
TypeScript
178 lines
6.4 KiB
TypeScript
'use client'
|
|
import React, { useCallback, useEffect, useState } from 'react'
|
|
import { useTranslation } from 'react-i18next'
|
|
import AppUnavailable from '../../base/app-unavailable'
|
|
import { ModelTypeEnum } from '../../header/account-setting/model-provider-page/declarations'
|
|
import StepOne from './step-one'
|
|
import StepTwo from './step-two'
|
|
import StepThree from './step-three'
|
|
import { Topbar } from './top-bar'
|
|
import { DataSourceType } from '@/models/datasets'
|
|
import type { CrawlOptions, CrawlResultItem, DataSet, FileItem, createDocumentResponse } from '@/models/datasets'
|
|
import { fetchDataSource } from '@/service/common'
|
|
import { fetchDatasetDetail } from '@/service/datasets'
|
|
import { DataSourceProvider, type NotionPage } from '@/models/common'
|
|
import { useModalContext } from '@/context/modal-context'
|
|
import { useDefaultModel } from '@/app/components/header/account-setting/model-provider-page/hooks'
|
|
|
|
type DatasetUpdateFormProps = {
|
|
datasetId?: string
|
|
}
|
|
|
|
const DEFAULT_CRAWL_OPTIONS: CrawlOptions = {
|
|
crawl_sub_pages: true,
|
|
only_main_content: true,
|
|
includes: '',
|
|
excludes: '',
|
|
limit: 10,
|
|
max_depth: '',
|
|
use_sitemap: true,
|
|
}
|
|
|
|
const DatasetUpdateForm = ({ datasetId }: DatasetUpdateFormProps) => {
|
|
const { t } = useTranslation()
|
|
const { setShowAccountSettingModal } = useModalContext()
|
|
const [hasConnection, setHasConnection] = useState(true)
|
|
const [dataSourceType, setDataSourceType] = useState<DataSourceType>(DataSourceType.FILE)
|
|
const [step, setStep] = useState(1)
|
|
const [indexingTypeCache, setIndexTypeCache] = useState('')
|
|
const [retrievalMethodCache, setRetrievalMethodCache] = useState('')
|
|
const [fileList, setFiles] = useState<FileItem[]>([])
|
|
const [result, setResult] = useState<createDocumentResponse | undefined>()
|
|
const [hasError, setHasError] = useState(false)
|
|
const { data: embeddingsDefaultModel } = useDefaultModel(ModelTypeEnum.textEmbedding)
|
|
|
|
const [notionPages, setNotionPages] = useState<NotionPage[]>([])
|
|
const updateNotionPages = (value: NotionPage[]) => {
|
|
setNotionPages(value)
|
|
}
|
|
|
|
const [websitePages, setWebsitePages] = useState<CrawlResultItem[]>([])
|
|
const [crawlOptions, setCrawlOptions] = useState<CrawlOptions>(DEFAULT_CRAWL_OPTIONS)
|
|
|
|
const updateFileList = (preparedFiles: FileItem[]) => {
|
|
setFiles(preparedFiles)
|
|
}
|
|
const [websiteCrawlProvider, setWebsiteCrawlProvider] = useState<DataSourceProvider>(DataSourceProvider.fireCrawl)
|
|
const [websiteCrawlJobId, setWebsiteCrawlJobId] = useState('')
|
|
|
|
const updateFile = (fileItem: FileItem, progress: number, list: FileItem[]) => {
|
|
const targetIndex = list.findIndex(file => file.fileID === fileItem.fileID)
|
|
list[targetIndex] = {
|
|
...list[targetIndex],
|
|
progress,
|
|
}
|
|
setFiles([...list])
|
|
// use follow code would cause dirty list update problem
|
|
// const newList = list.map((file) => {
|
|
// if (file.fileID === fileItem.fileID) {
|
|
// return {
|
|
// ...fileItem,
|
|
// progress,
|
|
// }
|
|
// }
|
|
// return file
|
|
// })
|
|
// setFiles(newList)
|
|
}
|
|
const updateIndexingTypeCache = (type: string) => {
|
|
setIndexTypeCache(type)
|
|
}
|
|
const updateResultCache = (res?: createDocumentResponse) => {
|
|
setResult(res)
|
|
}
|
|
const updateRetrievalMethodCache = (method: string) => {
|
|
setRetrievalMethodCache(method)
|
|
}
|
|
|
|
const nextStep = useCallback(() => {
|
|
setStep(step + 1)
|
|
}, [step, setStep])
|
|
|
|
const changeStep = useCallback((delta: number) => {
|
|
setStep(step + delta)
|
|
}, [step, setStep])
|
|
|
|
const checkNotionConnection = async () => {
|
|
const { data } = await fetchDataSource({ url: '/data-source/integrates' })
|
|
const hasConnection = data.filter(item => item.provider === 'notion') || []
|
|
setHasConnection(hasConnection.length > 0)
|
|
}
|
|
|
|
useEffect(() => {
|
|
checkNotionConnection()
|
|
}, [])
|
|
|
|
const [detail, setDetail] = useState<DataSet | null>(null)
|
|
useEffect(() => {
|
|
(async () => {
|
|
if (datasetId) {
|
|
try {
|
|
const detail = await fetchDatasetDetail(datasetId)
|
|
setDetail(detail)
|
|
}
|
|
catch (e) {
|
|
setHasError(true)
|
|
}
|
|
}
|
|
})()
|
|
}, [datasetId])
|
|
|
|
if (hasError)
|
|
return <AppUnavailable code={500} unknownReason={t('datasetCreation.error.unavailable') as string} />
|
|
|
|
return (
|
|
<div className='flex flex-col bg-components-panel-bg' style={{ height: 'calc(100vh - 56px)' }}>
|
|
<Topbar activeIndex={step - 1} />
|
|
<div style={{ height: 'calc(100% - 52px)' }}>
|
|
{step === 1 && <StepOne
|
|
hasConnection={hasConnection}
|
|
onSetting={() => setShowAccountSettingModal({ payload: 'data-source' })}
|
|
datasetId={datasetId}
|
|
dataSourceType={dataSourceType}
|
|
dataSourceTypeDisable={!!detail?.data_source_type}
|
|
changeType={setDataSourceType}
|
|
files={fileList}
|
|
updateFile={updateFile}
|
|
updateFileList={updateFileList}
|
|
notionPages={notionPages}
|
|
updateNotionPages={updateNotionPages}
|
|
onStepChange={nextStep}
|
|
websitePages={websitePages}
|
|
updateWebsitePages={setWebsitePages}
|
|
onWebsiteCrawlProviderChange={setWebsiteCrawlProvider}
|
|
onWebsiteCrawlJobIdChange={setWebsiteCrawlJobId}
|
|
crawlOptions={crawlOptions}
|
|
onCrawlOptionsChange={setCrawlOptions}
|
|
/>}
|
|
{(step === 2 && (!datasetId || (datasetId && !!detail))) && <StepTwo
|
|
isAPIKeySet={!!embeddingsDefaultModel}
|
|
onSetting={() => setShowAccountSettingModal({ payload: 'provider' })}
|
|
indexingType={detail?.indexing_technique}
|
|
datasetId={datasetId}
|
|
dataSourceType={dataSourceType}
|
|
files={fileList.map(file => file.file)}
|
|
notionPages={notionPages}
|
|
websitePages={websitePages}
|
|
websiteCrawlProvider={websiteCrawlProvider}
|
|
websiteCrawlJobId={websiteCrawlJobId}
|
|
onStepChange={changeStep}
|
|
updateIndexingTypeCache={updateIndexingTypeCache}
|
|
updateRetrievalMethodCache={updateRetrievalMethodCache}
|
|
updateResultCache={updateResultCache}
|
|
crawlOptions={crawlOptions}
|
|
/>}
|
|
{step === 3 && <StepThree
|
|
datasetId={datasetId}
|
|
datasetName={detail?.name}
|
|
indexingType={detail?.indexing_technique || indexingTypeCache}
|
|
retrievalMethod={detail?.retrieval_model_dict?.search_method || retrievalMethodCache}
|
|
creationCache={result}
|
|
/>}
|
|
</div>
|
|
</div>
|
|
)
|
|
}
|
|
|
|
export default DatasetUpdateForm
|